@misc{indiciaed5eab3781dff, title = {Align-Then-stEer: Adapting the Vision-Language Action Models through Unified Latent Guidance}, author = {Yang Zhang and Chenwei Wang and Ouyang Lu and Yuan Zhao and Yunfei Ge and Zhenglong Sun and Xiu Li and Chi Zhang and Chenjia Bai and Xuelong Li}, year = {2025}, url = {https://arxiv.org/abs/2509.02055}, note = {Source identifier: 2509.02055} }