@misc{indiciae93cfec038a4c, title = {Why Vision Language Models Struggle with Visual Arithmetic? Towards Enhanced Chart and Geometry Understanding}, author = {Kung-Hsiang Huang and Can Qin and Haoyi Qiu and Philippe Laban and Shafiq Joty and Caiming Xiong and Chien-Sheng Wu}, year = {2025}, url = {https://arxiv.org/abs/2502.11492}, note = {Source identifier: 2502.11492} }