@misc{indiciaeb4a34305c8fa, title = {VLM-3D:End-to-End Vision-Language Models for Open-World 3D Perception}, author = {Fuhao Chang and Shuxin Li and Yabei Li and Lei He}, year = {2025}, url = {https://arxiv.org/abs/2508.09061}, note = {Source identifier: 2508.09061} }