@misc{indiciae823b1f84489f, title = {Learning Compact Vision Tokens for Efficient Large Multimodal Models}, author = {Hao Tang and Chengchao Shen}, year = {2025}, url = {https://arxiv.org/abs/2506.07138}, note = {Source identifier: 2506.07138} }