@misc{indiciaefed00c43ce0d, title = {Visual Commonsense in Pretrained Unimodal and Multimodal Models}, author = {Chenyu Zhang and Benjamin Van Durme and Zhuowan Li and Elias Stengel-Eskin}, year = {2022}, url = {https://arxiv.org/abs/2205.01850}, note = {Source identifier: 2205.01850} }