@misc{indiciaee0c299b7020d, title = {Diffusion-Based Co-Speech Gesture Generation Using Joint Text and Audio Representation}, author = {Anna Deichler and Shivam Mehta and Simon Alexanderson and Jonas Beskow}, year = {2023}, doi = {10.1145/3577190.3616117}, url = {https://arxiv.org/abs/2309.05455}, note = {Source identifier: 2309.05455} }