@misc{indiciaed1bad0effbd2, title = {MASH-VLM: Mitigating Action-Scene Hallucination in Video-LLMs through Disentangled Spatial-Temporal Representations}, author = {Kyungho Bae and Jinhyung Kim and Sihaeng Lee and Soonyoung Lee and Gunhee Lee and Jinwoo Choi}, year = {2025}, url = {https://arxiv.org/abs/2503.15871}, note = {Source identifier: 2503.15871} }