@misc{indiciae0925db4de391, title = {Paparazzi: A Deep Dive into the Capabilities of Language and Vision Models for Grounding Viewpoint Descriptions}, author = {Henrik Voigt and Jan Hombeck and Monique Meuschke and Kai Lawonn and Sina Zarrieß}, year = {2023}, url = {https://arxiv.org/abs/2302.10282}, note = {Source identifier: 2302.10282} }