@misc{indiciae907267c63de1, title = {Contrastive Region Guidance: Improving Grounding in Vision-Language Models without Training}, author = {David Wan and Jaemin Cho and Elias Stengel-Eskin and Mohit Bansal}, year = {2024}, url = {https://arxiv.org/abs/2403.02325}, note = {Source identifier: 2403.02325} }