@misc{indiciae470e0b6e9a50, title = {Representation Before Training: A Practical Benchmark for Generative Medical Event Model Tokenization}, author = {Inhyeok Lee and Luke Solo and Michael C. Burkhart and Bashar Ramadan and Sahil Sethi and Sarah Jabbour and William F. Parker and Brett K. Beaulieu-Jones}, year = {2026}, url = {https://arxiv.org/abs/2604.16775}, note = {Source identifier: 2604.16775} }