@misc{indiciaecbaf193de805, title = {Keypoint-Integrated Instruction-Following Data Generation for Enhanced Human Pose and Action Understanding in Multimodal Models}, author = {Dewen Zhang and Wangpeng An and Hayaru Shouno}, year = {2025}, url = {https://arxiv.org/abs/2409.09306}, note = {Source identifier: 2409.09306} }