@misc{indiciae5804f9641960, title = {Unified Cross-modal Translation of Score Images, Symbolic Music, and Performance Audio}, author = {Jongmin Jung and Dongmin Kim and Sihun Lee and Seola Cho and Hyungjoon Soh and Irmak Bukey and Chris Donahue and Dasaem Jeong}, year = {2025}, doi = {10.1109/taslpro.2025.3648794}, url = {https://arxiv.org/abs/2505.12863}, note = {Source identifier: 2505.12863} }