@misc{indiciae49ff44364f98, title = {UNIMO-3: Multi-granularity Interaction for Vision-Language Representation Learning}, author = {Hao Yang and Can Gao and Hao Líu and Xinyan Xiao and Yanyan Zhao and Bing Qin}, year = {2023}, url = {https://arxiv.org/abs/2305.13697}, note = {Source identifier: 2305.13697} }