@misc{indiciae0898f58e477e, title = {Image-text Retrieval via Preserving Main Semantics of Vision}, author = {Xu Zhang and Xinzheng Niu and Philippe Fournier-Viger and Xudong Dai}, year = {2023}, url = {https://arxiv.org/abs/2304.10254}, note = {Source identifier: 2304.10254} }