@misc{indiciae5909ac3dc19b, title = {jina-vlm: Small Multilingual Vision Language Model}, author = {Andreas Koukounas and Georgios Mastrapas and Florian Hönicke and Sedigheh Eslami and Guillaume Roncari and Han Xiao}, year = {2026}, url = {https://arxiv.org/abs/2512.04032}, note = {Source identifier: 2512.04032} }