@misc{indiciae3c7973a4623a, title = {Caption-Driven Explorations: Aligning Image and Text Embeddings through Human-Inspired Foveated Vision}, author = {Dario Zanca and Andrea Zugarini and Simon Dietz and Thomas R. Altstidl and Mark A. Turban Ndjeuha and Leo Schwinn and Bjoern Eskofier}, year = {2024}, url = {https://arxiv.org/abs/2408.09948}, note = {Source identifier: 2408.09948} }