@misc{indiciae4ee986b5f35e, title = {cViL: Cross-Lingual Training of Vision-Language Models using Knowledge Distillation}, author = {Kshitij Gupta and Devansh Gautam and Radhika Mamidi}, year = {2022}, url = {https://arxiv.org/abs/2206.03354}, note = {Source identifier: 2206.03354} }