@misc{indiciae2e787fa81c0d, title = {Exploring Vision Transformers for 3D Human Motion-Language Models with Motion Patches}, author = {Qing Yu and Mikihiro Tanaka and Kent Fujiwara}, year = {2024}, url = {https://arxiv.org/abs/2405.04771}, note = {Source identifier: 2405.04771} }