@misc{indiciae5a52dd0eebd4, title = {See the Forest for the Trees: Loosely Speculative Decoding via Visual-Semantic Guidance for Efficient Inference of Video LLMs}, author = {Yicheng Ji and Jun Zhang and Jinpeng Chen and Cong Wang and Lidan Shou and Gang Chen and Huan Li}, year = {2026}, url = {https://arxiv.org/abs/2604.05650}, note = {Source identifier: 2604.05650} }