@misc{indiciae6435e2fb7e93, title = {EntityCLIP: Entity-Centric Image-Text Matching via Multimodal Attentive Contrastive Learning}, author = {Yaxiong Wang and Yujiao Wu and Lianwei Wu and Lechao Cheng and Zhun Zhong and Meng Wang}, year = {2025}, url = {https://arxiv.org/abs/2410.17810}, note = {Source identifier: 2410.17810} }