@misc{indiciaef5ffd4e3c6a4, title = {Structured Multi-modal Feature Embedding and Alignment for Image-Sentence Retrieval}, author = {Xuri Ge and Fuhai Chen and Joemon M. Jose and Zhilong Ji and Zhongqin Wu and Xiao Liu}, year = {2021}, doi = {10.1145/3474085.3475634}, url = {https://arxiv.org/abs/2108.02417}, note = {Source identifier: 2108.02417} }