@misc{indiciae2066115a3c15, title = {Informative Text-Image Alignment for Visual Affordance Learning with Foundation Models}, author = {Qian Zhang and Lin Zhang and Xing Fang and Mingxin Zhang and Zhiyuan Wei and Ran Song and Wei Zhang}, year = {2025}, url = {https://arxiv.org/abs/2509.17074}, note = {Source identifier: 2509.17074} }