@misc{indiciae566e3992d48e, title = {Textual Supervision Enhances Geospatial Representations in Vision-Language Models}, author = {Marcelo Sartori Locatelli and Fernando Tonucci and Jea Kwon and Luiz Felipe Vecchietti and Bryan Nathanael Wijaya and Cheng Yaw Low and Virgilio Almeida and Meeyoung Cha}, year = {2026}, url = {https://arxiv.org/abs/2606.07172}, note = {Source identifier: 2606.07172} }