@inproceedings{hwang-etal-2026-perceptual,
title = "Perceptual Hallucination in Vision{--}Language Models: Definition, Analysis and Verification",
author = "Hwang, Taewook and
Heo, Inbum and
Lee, Sung Jun and
Jung, Sangkeun",
editor = "Liakata, Maria and
Moreira, Viviane P. and
Zhang, Jiajun and
Jurgens, David",
booktitle = "Findings of the {A}ssociation for {C}omputational {L}inguistics: {ACL} 2026",
month = jul,
year = "2026",
address = "San Diego, California, United States",
publisher = "Association for Computational Linguistics",
url = "https://preview.aclanthology.org/ingest-acl/2026.findings-acl.1237/",
pages = "24710--24725",
ISBN = "979-8-89176-395-1",
abstract = "Vision-Language Models (VLMs) have demonstrated remarkable performance in document understanding tasks; however, VLMs also suffer from hallucinations inherited from LLMs. While prior work has focused on reasoning-stage hallucinations, the role of visual perception remains underexplored. In this work, we define perceptual hallucination as the phenomenon where VLMs generate information as if perceived, despite absent or damaged visual evidence. To analyze this, we construct DocHallu, a benchmark of 2,671 original{--}damaged image pairs across three tasks, available at https://huggingface.co/datasets/IB99/DocHallu. Experiments reveal that perceptual hallucination occurs across all models, with higher rates for numerical content than textual content. Activation patching analysis suggests that hallucinations are strongly associated with errors introduced in the vision encoder, which can subsequently propagate and become amplified through the text decoding process. We also demonstrate that LLM-based post-hoc filtering can reduce hallucination exposure by 36{\%} on average, with reductions of up to 88{\%}. This work extends VLM hallucination research by defining, analyzing, and verifying perceptual hallucination in document understanding."
}Markdown (Informal)
[Perceptual Hallucination in Vision–Language Models: Definition, Analysis and Verification](https://preview.aclanthology.org/ingest-acl/2026.findings-acl.1237/) (Hwang et al., Findings 2026)
ACL