@misc{indiciae9c167a2f2c90, title = {Semantics-Aware Human Motion Generation from Audio Instructions}, author = {Zi-An Wang and Shihao Zou and Shiyao Yu and Mingyuan Zhang and Chao Dong}, year = {2025}, doi = {10.1016/j.gmod.2025.101268}, url = {https://arxiv.org/abs/2505.23465}, note = {Source identifier: 2505.23465} }