@misc{indiciae806dae4b4e48, title = {VFaith: Do Large Multimodal Models Really Reason on Seen Images Rather than Previous Memories?}, author = {Jiachen Yu and Yufei Zhan and Ziheng Wu and Yousong Zhu and Jinqiao Wang and Minghui Qiu}, year = {2025}, url = {https://arxiv.org/abs/2506.11571}, note = {Source identifier: 2506.11571} }