@misc{indiciae919a04b25f08, title = {CLIP Models are Few-shot Learners: Empirical Studies on VQA and Visual Entailment}, author = {Haoyu Song and Li Dong and Wei-Nan Zhang and Ting Liu and Furu Wei}, year = {2022}, url = {https://arxiv.org/abs/2203.07190}, note = {Source identifier: 2203.07190} }