@inproceedings{smith-etal-2012-good,
title = "A good space: Lexical predictors in word space evaluation",
author = {Smith, Christian and
Danielsson, Henrik and
J{\"o}nsson, Arne},
editor = "Calzolari, Nicoletta and
Choukri, Khalid and
Declerck, Thierry and
Do{\u{g}}an, Mehmet U{\u{g}}ur and
Maegaard, Bente and
Mariani, Joseph and
Moreno, Asuncion and
Odijk, Jan and
Piperidis, Stelios",
booktitle = "Proceedings of the Eighth International Conference on Language Resources and Evaluation ({LREC}`12)",
month = may,
year = "2012",
address = "Istanbul, Turkey",
publisher = "European Language Resources Association (ELRA)",
url = "https://preview.aclanthology.org/add-emnlp-2024-awards/L12-1159/",
pages = "2530--2535",
abstract = "Vector space models benefit from using an outside corpus to train the model. It is, however, unclear what constitutes a good training corpus. We have investigated the effect on summary quality when using various language resources to train a vector space based extraction summarizer. This is done by evaluating the performance of the summarizer utilizing vector spaces built from corpora from different genres, partitioned from the Swedish SUC-corpus. The corpora are also characterized using a variety of lexical measures commonly used in readability studies. The performance of the summarizer is measured by comparing automatically produced summaries to human created gold standard summaries using the ROUGE F-score. Our results show that the genre of the training corpus does not have a significant effect on summary quality. However, evaluating the variance in the F-score between the genres based on lexical measures as independent variables in a linear regression model, shows that vector spaces created from texts with high syntactic complexity, high word variation, short sentences and few long words produce better summaries."
}
Markdown (Informal)
[A good space: Lexical predictors in word space evaluation](https://preview.aclanthology.org/add-emnlp-2024-awards/L12-1159/) (Smith et al., LREC 2012)
ACL
- Christian Smith, Henrik Danielsson, and Arne Jönsson. 2012. A good space: Lexical predictors in word space evaluation. In Proceedings of the Eighth International Conference on Language Resources and Evaluation (LREC'12), pages 2530–2535, Istanbul, Turkey. European Language Resources Association (ELRA).