@misc{indiciae5515d0f3aea3, title = {SparseOccVLA: Bridging Occupancy and Vision-Language Models via Sparse Queries for Unified 4D Scene Understanding and Planning}, author = {Chenxu Dang and Jie Wang and Guang Li and Zhiwen Hou and Zihan You and Hangjun Ye and Jie Ma and Long Chen and Yan Wang}, year = {2026}, url = {https://arxiv.org/abs/2601.06474}, note = {Source identifier: 2601.06474} }