@misc{indiciaedf5824d558a4, title = {Generate to Ground: Multimodal Text Conditioning Boosts Phrase Grounding in Medical Vision-Language Models}, author = {Felix Nützel and Mischa Dombrowski and Bernhard Kainz}, year = {2025}, url = {https://arxiv.org/abs/2507.12236}, note = {Source identifier: 2507.12236} }