@inproceedings{singh-etal-2023-many,
title = "Too Many Cooks Spoil the Model: Are Bilingual Models for {S}lovene Better than a Large Multilingual Model?",
author = "Singh, Pranaydeep and
Maladry, Aaron and
Lefever, Els",
editor = "Piskorski, Jakub and
Marci{\'n}czuk, Micha{\l} and
Nakov, Preslav and
Ogrodniczuk, Maciej and
Pollak, Senja and
P{\v{r}}ib{\'a}{\v{n}}, Pavel and
Rybak, Piotr and
Steinberger, Josef and
Yangarber, Roman",
booktitle = "Proceedings of the 9th Workshop on Slavic Natural Language Processing 2023 (SlavicNLP 2023)",
month = may,
year = "2023",
address = "Dubrovnik, Croatia",
publisher = "Association for Computational Linguistics",
url = "https://preview.aclanthology.org/add-emnlp-2024-awards/2023.bsnlp-1.5/",
doi = "10.18653/v1/2023.bsnlp-1.5",
pages = "32--39",
abstract = "This paper investigates whether adding data of typologically closer languages improves the performance of transformer-based models for three different downstream tasks, namely Part-of-Speech tagging, Named Entity Recognition, and Sentiment Analysis, compared to a monolingual and plain multilingual language model. For the presented pilot study, we performed experiments for the use case of Slovene, a low(er)-resourced language belonging to the Slavic language family. The experiments were carried out in a controlled setting, where a monolingual model for Slovene was compared to combined language models containing Slovene, trained with the same amount of Slovene data. The experimental results show that adding typologically closer languages indeed improves the performance of the Slovene language model, and even succeeds in outperforming the large multilingual XLM-RoBERTa model for NER and PoS-tagging. We also reveal that, contrary to intuition, distantly or unrelated languages also combine admirably with Slovene, often out-performing XLM-R as well. All the bilingual models used in the experiments are publicly available at \url{https://github.com/pranaydeeps/BLAIR}"
}
Markdown (Informal)
[Too Many Cooks Spoil the Model: Are Bilingual Models for Slovene Better than a Large Multilingual Model?](https://preview.aclanthology.org/add-emnlp-2024-awards/2023.bsnlp-1.5/) (Singh et al., BSNLP 2023)
ACL