@misc{indiciaeb3ade830fd51, title = {Training Vision-Language-Action Models with Dense Embodied Chain-of-Thought Supervision}, author = {Haoyang Li and Guanlin Li and Youhe Feng and Chen Zhao and Zhuoran Wang and Yang Li and Qizhe Wei and Shifeng Bao and Haitao Shen and Yihan Zhao and Tong Yang and Jing Zhang}, year = {2026}, url = {https://arxiv.org/abs/2606.30552}, note = {Source identifier: 2606.30552} }