Closing the Feedback Loop: From Experience Extraction to Insight Governance in Verbal Reinforcement Learning
- 2026-06-16
- 4
@misc{cui2026closing,
title = {Closing the Feedback Loop: From Experience Extraction to Insight Governance in Verbal Reinforcement Learning},
author = {Ya Cui and Xing Zhang and Yulong Zhang and Lingzhi Shao and Xiaofeng Shi and Guanghui Wang and Pei-Gen He},
year = {2026},
eprint = {2606.17591},
archivePrefix = {arXiv},
url = {https://arxiv.org/abs/2606.17591},
}