@misc{indiciaef0608733690c, title = {What do vision-language models see in the context? Investigating multimodal in-context learning}, author = {Gabriel O. dos Santos and Esther Colombini and Sandra Avila}, year = {2025}, url = {https://arxiv.org/abs/2510.24331}, note = {Source identifier: 2510.24331} }