@inproceedings{masis-etal-2022-corpus,
title = "Corpus-Guided Contrast Sets for Morphosyntactic Feature Detection in Low-Resource {E}nglish Varieties",
author = "Masis, Tessa and
Neal, Anissa and
Green, Lisa and
O{'}Connor, Brendan",
booktitle = "Proceedings of the first workshop on NLP applications to field linguistics",
month = oct,
year = "2022",
address = "Gyeongju, Republic of Korea",
publisher = "International Conference on Computational Linguistics",
url = "https://aclanthology.org/2022.fieldmatters-1.2",
pages = "11--25",
abstract = "The study of language variation examines how language varies between and within different groups of speakers, shedding light on how we use language to construct identities and how social contexts affect language use. A common method is to identify instances of a certain linguistic feature - say, the zero copula construction - in a corpus, and analyze the feature{'}s distribution across speakers, topics, and other variables, to either gain a qualitative understanding of the feature{'}s function or systematically measure variation. In this paper, we explore the challenging task of automatic morphosyntactic feature detection in low-resource English varieties. We present a human-in-the-loop approach to generate and filter effective contrast sets via corpus-guided edits. We show that our approach improves feature detection for both Indian English and African American English, demonstrate how it can assist linguistic research, and release our fine-tuned models for use by other researchers.",
}
Markdown (Informal)
[Corpus-Guided Contrast Sets for Morphosyntactic Feature Detection in Low-Resource English Varieties](https://aclanthology.org/2022.fieldmatters-1.2) (Masis et al., FieldMatters 2022)
ACL