@misc{indiciaec4a3fff5531b, title = {RL\$\textasciicircum{}2\$-VLA: Adaptive RL Latent Compositional Steering with Test-Time Scaling for Vision-Language-Action Models}, author = {Derek Ming Siang Tan and Shailesh Shailesh and Srikrishna Iyer and William Wei Jie Teo and Yuanliang Ju and Qiao Gu and Guillaume Sartoretti}, year = {2026}, url = {https://arxiv.org/abs/2607.26991}, note = {Source identifier: 2607.26991} }