@misc{indiciae057ea13a920a, title = {Your Large Vision-Language Model Only Needs A Few Attention Heads For Visual Grounding}, author = {Seil Kang and Jinyeong Kim and Junhyeok Kim and Seong Jae Hwang}, year = {2025}, url = {https://arxiv.org/abs/2503.06287}, note = {Source identifier: 2503.06287} }