@misc{indiciae781fe92b00ae, title = {KG-ViP: Bridging Knowledge Grounding and Visual Perception in Multi-modal LLMs for Visual Question Answering}, author = {Zhiyang Li and Ao Ke and Yukun Cao and Xike Xie}, year = {2026}, url = {https://arxiv.org/abs/2601.11632}, note = {Source identifier: 2601.11632} }