@misc{indiciae1f1019449c80, title = {Distilling 3D Spatial Reasoning into a Lightweight Vision-Language Model with CoT}, author = {Alaa Asfour and Christopher Indris and Leihan Chen and Tejas Vyas and Guanghui Wang}, year = {2026}, url = {https://arxiv.org/abs/2605.09719}, note = {Source identifier: 2605.09719} }