@misc{indiciaec922827eb69f, title = {Training Vision-Language Models with Less Bimodal Supervision}, author = {Elad Segal and Ben Bogin and Jonathan Berant}, year = {2022}, url = {https://arxiv.org/abs/2211.00262}, note = {Source identifier: 2211.00262} }