@misc{indiciae1affeaeb4c40, title = {Video-to-Audio Generation with Fine-grained Temporal Semantics}, author = {Yuchen Hu and Yu Gu and Chenxing Li and Rilin Chen and Dong Yu}, year = {2024}, url = {https://arxiv.org/abs/2409.14709}, note = {Source identifier: 2409.14709} }