@misc{indiciae6753bf5e1073, title = {Visual cognition in multimodal large language models}, author = {Luca M. Schulze Buschoff and Elif Akata and Matthias Bethge and Eric Schulz}, year = {2024}, url = {https://arxiv.org/abs/2311.16093}, note = {Source identifier: 2311.16093} }