@misc{indiciaecebd5dc7a593, title = {Deep Audio-Visual Singing Voice Transcription based on Self-Supervised Learning Models}, author = {Xiangming Gu and Wei Zeng and Jianan Zhang and Longshen Ou and Ye Wang}, year = {2023}, url = {https://arxiv.org/abs/2304.12082}, note = {Source identifier: 2304.12082} }