@misc{indiciaeb6fcedc821b5, title = {Token-level Response-visual Attention Guidance for Multimodal LLMs Knowledge Distillation}, author = {Jaehyun Jang and Eunseop Yoon and Hee Suk Yoon and SooHwan Eom and Mark A. Hasegawa-Johnson and Chang D. Yoo}, year = {2026}, url = {https://arxiv.org/abs/2607.02593}, note = {Source identifier: 2607.02593} }