@inproceedings{rucker-akbik-2025-evaluating,
title = "Evaluating Design Decisions for Dual Encoder-based Entity Disambiguation",
author = {R{\"u}cker, Susanna and
Akbik, Alan},
editor = "Che, Wanxiang and
Nabende, Joyce and
Shutova, Ekaterina and
Pilehvar, Mohammad Taher",
booktitle = "Proceedings of the 63rd Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers)",
month = jul,
year = "2025",
address = "Vienna, Austria",
publisher = "Association for Computational Linguistics",
url = "https://preview.aclanthology.org/ingestion-acl-25/2025.acl-long.764/",
pages = "15685--15701",
ISBN = "979-8-89176-251-0",
abstract = "Entity disambiguation (ED) is the task of linking mentions in text to corresponding entries in a knowledge base. Dual Encoders address this by embedding mentions and label candidates in a shared embedding space and applying a similarity metric to predict the correct label. In this work, we focus on evaluating key design decisions for Dual Encoder-based ED, such as its loss function, similarity metric, label verbalization format, and negative sampling strategy. We present the resulting model VerbalizED, a document-level Dual Encoder model that includes contextual label verbalizations and efficient hard negative sampling. Additionally, we explore an iterative prediction variant that aims to improve the disambiguation of challenging data points. To support our analysis, we first conduct comprehensive ablation experiments on specific design decisions using AIDA-Yago, followed by large-scale, multi-domain evaluation on the ZELDA benchmark."
}
Markdown (Informal)
[Evaluating Design Decisions for Dual Encoder-based Entity Disambiguation](https://preview.aclanthology.org/ingestion-acl-25/2025.acl-long.764/) (Rücker & Akbik, ACL 2025)
ACL