@misc{indiciae3122e424a370, title = {Pointing-VLA: Typed Spatial Grounding Interfaces for Vision-Language-Action Manipulation}, author = {Xiwen Chen and Zelin Li and Zhiruo Zhou and Huiming Chen and Chenwei Wang and Xiaojun Zhu}, year = {2026}, url = {https://arxiv.org/abs/2608.23138}, note = {Source identifier: 2608.23138} }