@misc{indiciae7a1259120d27, title = {STA-V2A: Video-to-Audio Generation with Semantic and Temporal Alignment}, author = {Yong Ren and Chenxing Li and Manjie Xu and Wei Liang and Yu Gu and Rilin Chen and Dong Yu}, year = {2025}, url = {https://arxiv.org/abs/2409.08601}, note = {Source identifier: 2409.08601} }