@misc{indiciae5dea8b55027e, title = {See then Tell: Enhancing Key Information Extraction with Vision Grounding}, author = {Shuhang Liu and Zhenrong Zhang and Pengfei Hu and Jiefeng Ma and Jun Du and Qing Wang and Jianshu Zhang and Chenyu Liu}, year = {2025}, url = {https://arxiv.org/abs/2409.19573}, note = {Source identifier: 2409.19573} }