@inproceedings{dfc58ac538d845bcad55d4e5788070da,
title = "Automatic discovery and transfer of MAXQ hierarchies in a complex system",
abstract = "Reinforcement learning has been an important category of machine learning approaches exhibiting self-learning and online learning characteristics. Using reinforcement learning, an agent can learn its behaviors through trial-and-error interactions with a dynamic environment and finally come up with an optimal strategy. Reinforcement learning suffers the curse of dimensionality, though there has been significant progress to overcome this issue in recent years. MAXQ is one of the most common approaches for reinforcement learning. To function properly, MAXQ requires a decomposition of the agent's task into a task hierarchy. Previously, the decomposition can only be done manually. In this paper, we propose a mechanism for automatic subtask discovery. The mechanism applies clustering to automatically construct task hierarchy required by MAXQ, such that MAXQ can be fully automated. We present the design of our mechanism, and demonstrate its effectiveness through theoretical analysis and an extensive experimental evaluation.",
keywords = "Clustering, MAXQ, Reinforcement Learning, System of Systems",
author = "Hongbing Wang and Wenya Li and Xuan Zhou",
year = "2012",
doi = "10.1109/ICTAI.2012.165",
language = "英语",
isbn = "9780769549156",
series = "Proceedings - International Conference on Tools with Artificial Intelligence, ICTAI",
publisher = "IEEE Computer Society",
pages = "1157--1162",
booktitle = "Proceedings - 2012 IEEE 24th International Conference on Tools with Artificial Intelligence, ICTAI 2012",
address = "美国",
note = "24th IEEE International Conference on Tools with Artificial Intelligence, ICTAI 2012 ; Conference date: 07-11-2012 Through 09-11-2012",
}