@misc{indiciae432d3562bcf0, title = {Teaching Vision-Language Models to Use the Scale They Are Given: Label-Free Equivariance Training for Metric Physical Reasoning}, author = {Kaizhen Tan and Yang Feng and Heqing Du and Siru Tao and Xin Xu and Hanzhe Hong}, year = {2026}, url = {https://arxiv.org/abs/2609.00658}, note = {Source identifier: 2609.00658} }