@misc{indiciae40648e61bb98, title = {Instruction-Guided Fusion of Multi-Layer Visual Features in Large Vision-Language Models}, author = {Xu Li and Yi Zheng and Haotian Chen and Xiaolei Chen and Yuxuan Liang and Chenghang Lai and Bin Li and Xiangyang Xue}, year = {2025}, url = {https://arxiv.org/abs/2501.08443}, note = {Source identifier: 2501.08443} }