Just-In-Time Reinforcement Learning: Continual Learning in LLM Agents Without Gradient Updates
- 2026-01-26
- 11
@misc{li2026justintime,
title = {Just-In-Time Reinforcement Learning: Continual Learning in LLM Agents Without Gradient Updates},
author = {Yi-Bo Li and Zijie Lin and Ailin Deng and Xuan Zhang and Yu-Fei He and Shuo Ji and Tri Cao and Bryan Hooi},
year = {2026},
eprint = {2601.18510},
archivePrefix = {arXiv},
url = {https://arxiv.org/abs/2601.18510},
}