@inproceedings{sperduti-nguyen-2025-pset,
title = "{PSET}: a Phonetics-Semantics Evaluation Testbed",
author = "Sperduti, Gianluca and
Nguyen, Dong",
editor = "Christodoulopoulos, Christos and
Chakraborty, Tanmoy and
Rose, Carolyn and
Peng, Violet",
booktitle = "Proceedings of the 2025 Conference on Empirical Methods in Natural Language Processing",
month = nov,
year = "2025",
address = "Suzhou, China",
publisher = "Association for Computational Linguistics",
url = "https://preview.aclanthology.org/ingest-emnlp/2025.emnlp-main.373/",
pages = "7357--7367",
ISBN = "979-8-89176-332-6",
abstract = "We introduce the Phonetics-Semantics Evaluation Testbed (PSET), a new English-based testbed to evaluate phonetic embeddings. Our testbed is built on the assumption that phonetic embeddings should always prioritize phonetics over semantics, and it therefore leverages homophones and synonyms.We use PSET to test three phonetic embedding models: articulatory embeddings, Phoneme2Vec, and XPhoneBERT. The phonetic-based embeddings solve the task with varying degrees of success, with Phoneme2Vec performing the best.We also test five recent LLMs, GPT-4o, Gemini 2.5 Flash, Llama 3.1-8B, OLMo-7B and OLMo 2-7B. Gemini 2.5 Flash performs better than the other models. With this testbed, we hope to advance the development and evaluation of phonetic embedding models."
}Markdown (Informal)
[PSET: a Phonetics-Semantics Evaluation Testbed](https://preview.aclanthology.org/ingest-emnlp/2025.emnlp-main.373/) (Sperduti & Nguyen, EMNLP 2025)
ACL
- Gianluca Sperduti and Dong Nguyen. 2025. PSET: a Phonetics-Semantics Evaluation Testbed. In Proceedings of the 2025 Conference on Empirical Methods in Natural Language Processing, pages 7357–7367, Suzhou, China. Association for Computational Linguistics.