@misc{indiciae71633cc9b852, title = {Visual Anchors Are Strong Information Aggregators For Multimodal Large Language Model}, author = {Haogeng Liu and Quanzeng You and Xiaotian Han and Yongfei Liu and Huaibo Huang and Ran He and Hongxia Yang}, year = {2024}, url = {https://arxiv.org/abs/2405.17815}, note = {Source identifier: 2405.17815} }