@misc{indiciae72c37dd607c6, title = {XS-VLA: Teaching Tiny Vision-Language-Action Models with Spatial Supervision and Demonstration Conditioning}, author = {Iok Tong Lei and Ying Jie Yap and Wei Huang and Qingchen Xie and Qianzhi Li and Yujie Zhang and Xiaolong Liu and Zhidong Deng}, year = {2026}, url = {https://arxiv.org/abs/2607.04171}, note = {Source identifier: 2607.04171} }