@misc{indiciae568fe3600827, title = {Prototype-Aware Multimodal Alignment for Open-Vocabulary Visual Grounding}, author = {Jiangnan Xie and Xiaolong Zheng and Liang Zheng}, year = {2025}, url = {https://arxiv.org/abs/2509.06291}, note = {Source identifier: 2509.06291} }