@misc{indiciae530ae2ffbd9b, title = {3DThinkVLA: Endowing Vision-Language-Action Models with Latent 3D Priors via 3D-Thinking-Guided Co-training}, author = {Jiaxin Shi and Xidong Zhang and Fucai Zhu and Zhe Li and Siyu Zhu and Weihao Yuan}, year = {2026}, url = {https://arxiv.org/abs/2606.04436}, note = {Source identifier: 2606.04436} }