@misc{indiciaee9aec07a5533, title = {HandVQA: Diagnosing and Improving Fine-Grained Spatial Reasoning about Hands in Vision-Language Models}, author = {MD Khalequzzaman Chowdhury Sayem and Mubarrat Tajoar Chowdhury and Yihalem Yimolal Tiruneh and Muneeb A. Khan and Muhammad Salman Ali and Binod Bhattarai and Seungryul Baek}, year = {2026}, url = {https://arxiv.org/abs/2603.26362}, note = {Source identifier: 2603.26362} }