@misc{indiciaed66dd1c14ba5, title = {VideoAesBench: Benchmarking the Video Aesthetics Perception Capabilities of Large Multimodal Models}, author = {Yunhao Li and Sijing Wu and Zhilin Gao and Zicheng Zhang and Qi Jia and Huiyu Duan and Xiongkuo Min and Guangtao Zhai}, year = {2026}, url = {https://arxiv.org/abs/2601.21915}, note = {Source identifier: 2601.21915} }