@misc{indiciae8a3eebab6b59, title = {BridgeTower: Building Bridges Between Encoders in Vision-Language Representation Learning}, author = {Xiao Xu and Chenfei Wu and Shachar Rosenman and Vasudev Lal and Wanxiang Che and Nan Duan}, year = {2024}, url = {https://arxiv.org/abs/2206.08657}, note = {Source identifier: 2206.08657} }