@misc{indiciae3f23010c10e3, title = {LLaVA-SP: Enhancing Visual Representation with Visual Spatial Tokens for MLLMs}, author = {Haoran Lou and Chunxiao Fan and Ziyan Liu and Yuexin Wu and Xinliang Wang}, year = {2025}, url = {https://arxiv.org/abs/2507.00505}, note = {Source identifier: 2507.00505} }