@misc{indiciae08e042daf12e, title = {MultiToP: Learning to Patch Visual Tokens to Mitigate Hallucinations in Video Large Multimodal Models}, author = {Yuansheng Gao and Wenbin Xing and Jiahao Yuan and Kaiwen Zhou and Han Bao and Zonghui Wang and Wenzhi Chen}, year = {2026}, url = {https://arxiv.org/abs/2606.11792}, note = {Source identifier: 2606.11792} }