@inproceedings{rose-etal-2024-one,
title = "One-Vs-Rest Neural Network {E}nglish Grapheme Segmentation: A Linguistic Perspective",
author = "Rose, Samuel and
Dethlefs, Nina and
Kambhampati, C.",
editor = "Barak, Libby and
Alikhani, Malihe",
booktitle = "Proceedings of the 28th Conference on Computational Natural Language Learning",
month = nov,
year = "2024",
address = "Miami, FL, USA",
publisher = "Association for Computational Linguistics",
url = "https://preview.aclanthology.org/jlcl-multiple-ingestion/2024.conll-1.36/",
doi = "10.18653/v1/2024.conll-1.36",
pages = "464--469",
abstract = "Grapheme-to-Phoneme (G2P) correspondences form foundational frameworks of tasks such as text-to-speech (TTS) synthesis or automatic speech recognition. The G2P process involves taking words in their written form and generating their pronunciation. In this paper, we critique the status quo definition of a grapheme, currently a forced alignment process relating a single character to either a phoneme or a blank unit, that underlies the majority of modern approaches. We develop a linguistically-motivated redefinition from simple concepts such as vowel and consonant count and word length and offer a proof-of-concept implementation based on a multi-binary neural classification task. Our model achieves state-of-the-art results with a 31.86{\%} Word Error Rate on a standard benchmark, while generating linguistically meaningful grapheme segmentations."
}
Markdown (Informal)
[One-Vs-Rest Neural Network English Grapheme Segmentation: A Linguistic Perspective](https://preview.aclanthology.org/jlcl-multiple-ingestion/2024.conll-1.36/) (Rose et al., CoNLL 2024)
ACL