@misc{indiciae451c9439f142, title = {Text-Audio-Visual-conditioned Diffusion Model for Video Saliency Prediction}, author = {Li Yu and Xuanzhe Sun and Wei Zhou and Moncef Gabbouj}, year = {2025}, url = {https://arxiv.org/abs/2504.14267}, note = {Source identifier: 2504.14267} }