@inproceedings{patra-etal-2026-improving,
title = "Improving Language Identification for Code-Switched Speech: The Pivotal Role of Accented {E}nglish",
author = "Patra, Adyasha and
Sah, Dhiraj Kumar and
Jyothi, Preethi",
editor = "Demberg, Vera and
Inui, Kentaro and
Marquez, Llu{\'i}s",
booktitle = "Findings of the {A}ssociation for {C}omputational {L}inguistics: {EACL} 2026",
month = mar,
year = "2026",
address = "Rabat, Morocco",
publisher = "Association for Computational Linguistics",
url = "https://preview.aclanthology.org/ingest-eacl/2026.findings-eacl.242/",
pages = "4643--4656",
ISBN = "979-8-89176-386-9",
abstract = "Code-switching, where speakers alternate between languages within a single utterance, poses unique challenges for language identification (LID). Existing LID models often fail to reliably identify English spoken with the accent of the matrix (dominant) language. We show that finetuning LID models with small amounts of such accented English significantly improves code-switched LID, without degrading performance on standard monolingual speech{---}a limitation observed with direct finetuning on code-switched utterances. This is achieved via low-rank adaptation (LoRA) on limited accented data, which allows models to adapt efficiently. To better evaluate performance, we introduce LangRank, a metric that captures the relative ranking of identified languages often overlooked by traditional metrics. Our method generalizes across multiple language pairs, including Hindi-English, Bengali-English, Mandarin-English, and Arabic-English, providing robust LID in code-switched multilingual contexts."
}Markdown (Informal)
[Improving Language Identification for Code-Switched Speech: The Pivotal Role of Accented English](https://preview.aclanthology.org/ingest-eacl/2026.findings-eacl.242/) (Patra et al., Findings 2026)
ACL