@misc{indiciae935aa5a6c058, title = {LLaVA-4D: Embedding SpatioTemporal Prompt into LMMs for 4D Scene Understanding}, author = {Hanyu Zhou and Gim Hee Lee}, year = {2025}, url = {https://arxiv.org/abs/2505.12253}, note = {Source identifier: 2505.12253} }