@inproceedings{palma-gomez-rozovskaya-2024-multi,
title = "Multi-Reference Benchmarks for {R}ussian Grammatical Error Correction",
author = "Palma Gomez, Frank and
Rozovskaya, Alla",
editor = "Graham, Yvette and
Purver, Matthew",
booktitle = "Proceedings of the 18th Conference of the European Chapter of the Association for Computational Linguistics (Volume 1: Long Papers)",
month = mar,
year = "2024",
address = "St. Julian{'}s, Malta",
publisher = "Association for Computational Linguistics",
url = "https://preview.aclanthology.org/jlcl-multiple-ingestion/2024.eacl-long.76/",
pages = "1253--1270",
abstract = "This paper presents multi-reference benchmarks for the Grammatical Error Correction (GEC) of Russian, based on two existing single-reference datasets, for a total of 7,444 learner sentences from a variety of first language backgrounds. Each sentence is corrected independently by two new raters, and their corrections are reviewed by a senior annotator, resulting in a total of three references per sentence. Analysis of the annotations reveals that the new raters tend to make more changes, compared to the original raters, especially at the lexical level. We conduct experiments with two popular GEC approaches and show competitive performance on the original datasets and the new benchmarks. We also compare system scores as evaluated against individual annotators and discuss the effect of using multiple references overall and on specific error types. We find that using the union of the references increases system scores by more than 10 points and decreases the gap between system and human performance, thereby providing a more realistic evaluation of GEC system performance, although the effect is not the same across the error types. The annotations are available for research."
}
Markdown (Informal)
[Multi-Reference Benchmarks for Russian Grammatical Error Correction](https://preview.aclanthology.org/jlcl-multiple-ingestion/2024.eacl-long.76/) (Palma Gomez & Rozovskaya, EACL 2024)
ACL