@misc{indiciae85a4036e3740, title = {HumanVLM: Foundation for Human-Scene Vision-Language Model}, author = {Dawei Dai and Xu Long and Li Yutang and Zhang Yuanhui and Shuyin Xia}, year = {2024}, url = {https://arxiv.org/abs/2411.03034}, note = {Source identifier: 2411.03034} }