@misc{indiciae1fd78118105b, title = {What Vision-Language Models `See' when they See Scenes}, author = {Michele Cafagna and Kees van Deemter and Albert Gatt}, year = {2021}, url = {https://arxiv.org/abs/2109.07301}, note = {Source identifier: 2109.07301} }