@misc{indiciaec69f3599af52, title = {What Keeps Vision-Language Models Looking at the Image?}, author = {Hiroto Osaka and Shohei Taniguchi and Gouki Minegishi and Kai Yamashita and Masahiro Suzuki and Yutaka Matsuo}, year = {2026}, url = {https://arxiv.org/abs/2607.12815}, note = {Source identifier: 2607.12815} }