@misc{indiciae4c14ecddaf97, title = {ViSIL: Unified Evaluation of Information Loss in Multimodal Video Captioning}, author = {Po-han Li and Shenghui Chen and Ufuk Topcu and Sandeep Chinchali}, year = {2026}, url = {https://arxiv.org/abs/2601.09851}, note = {Source identifier: 2601.09851} }