@misc{indiciae11ce4421b210, title = {Infusing fine-grained visual knowledge to Vision-Language Models}, author = {Nikolaos-Antonios Ypsilantis and Kaifeng Chen and André Araujo and Ondřej Chum}, year = {2025}, url = {https://arxiv.org/abs/2508.12137}, note = {Source identifier: 2508.12137} }