@inproceedings{vasilyev-bohannon-2021-estime,
title = "{ESTIME}: Estimation of Summary-to-Text Inconsistency by Mismatched Embeddings",
author = "Vasilyev, Oleg and
Bohannon, John",
editor = "Gao, Yang and
Eger, Steffen and
Zhao, Wei and
Lertvittayakumjorn, Piyawat and
Fomicheva, Marina",
booktitle = "Proceedings of the 2nd Workshop on Evaluation and Comparison of NLP Systems",
month = nov,
year = "2021",
address = "Punta Cana, Dominican Republic",
publisher = "Association for Computational Linguistics",
url = "https://preview.aclanthology.org/landing_page/2021.eval4nlp-1.10/",
doi = "10.18653/v1/2021.eval4nlp-1.10",
pages = "94--103",
abstract = "We propose a new reference-free summary quality evaluation measure, with emphasis on the faithfulness. The measure is based on finding and counting all probable potential inconsistencies of the summary with respect to the source document. The proposed ESTIME, Estimator of Summary-to-Text Inconsistency by Mismatched Embeddings, correlates with expert scores in summary-level SummEval dataset stronger than other common evaluation measures not only in Consistency but also in Fluency. We also introduce a method of generating subtle factual errors in human summaries. We show that ESTIME is more sensitive to subtle errors than other common evaluation measures."
}
Markdown (Informal)
[ESTIME: Estimation of Summary-to-Text Inconsistency by Mismatched Embeddings](https://preview.aclanthology.org/landing_page/2021.eval4nlp-1.10/) (Vasilyev & Bohannon, Eval4NLP 2021)
ACL