@misc{indiciae62c49fca2b5d, title = {AlignCAT: Visual-Linguistic Alignment of Category and Attribute for Weakly Supervised Visual Grounding}, author = {Yidan Wang and Chenyi Zhuang and Wutao Liu and Pan Gao and Nicu Sebe}, year = {2025}, url = {https://arxiv.org/abs/2508.03201}, note = {Source identifier: 2508.03201} }