@inproceedings{38e916a77d964f3d8c9816a8d1a96a25,
title = "Integrating partial model knowledge in model free RL algorithms",
abstract = "In reinforcement learning an agent uses online feedback from the environment and prior knowledge in order to adaptively select an effective policy. Model free approaches address this task by directly mapping external and internal states to actions, while model based methods attempt to construct a model of the environment, followed by a selection of optimal actions based on that model. Given the complementary advantages of both approaches, we suggest a novel algorithm which combines them into a single algorithm, which switches between a model based and a model free mode, depending on the current environmental state and on the status of the agent's knowledge. We prove that such an approach leads to improved performance whenever environmental knowledge is available, without compromising performance when such knowledge is absent. Numerical simulations demonstrate the effectiveness of the approach and suggest its efficacy in boosting policy gradient learning.",
author = "Aviv Tamar and \{Di Castro\}, Dotan and Ron Meir",
year = "2011",
language = "الإنجليزيّة",
series = "Proceedings of the 28th International Conference on Machine Learning, ICML 2011",
pages = "305--312",
booktitle = "Proceedings of the 28th International Conference on Machine Learning, ICML 2011",
note = "28th International Conference on Machine Learning, ICML 2011 ; Conference date: 28-06-2011 Through 02-07-2011",
}