@misc{indiciae4e5736944461, title = {Fusion to Enhance: Fusion Visual Encoder to Enhance Multimodal Language Model}, author = {Yifei She and Huangxuan Wu}, year = {2025}, url = {https://arxiv.org/abs/2509.00664}, note = {Source identifier: 2509.00664} }