Stefan O'Toole, Nir Lipovetzky, Miquel Ramírez, Adrian R. Pearce. Width-based Lookaheads with Learnt Base Policies and Heuristics Over the Atari-2600 Benchmark. In Marc'Aurelio Ranzato, Alina Beygelzimer, Yann N. Dauphin, Percy Liang, Jennifer Wortman Vaughan, editors, Advances in Neural Information Processing Systems 34: Annual Conference on Neural Information Processing Systems 2021, NeurIPS 2021, December 6-14, 2021, virtual. pages 26536-26547, 2021. [doi]
@inproceedings{OTooleLRP21, title = {Width-based Lookaheads with Learnt Base Policies and Heuristics Over the Atari-2600 Benchmark}, author = {Stefan O'Toole and Nir Lipovetzky and Miquel Ramírez and Adrian R. Pearce}, year = {2021}, url = {https://proceedings.neurips.cc/paper/2021/hash/df42e2244c97a0d80d565ae8176d3351-Abstract.html}, researchr = {https://researchr.org/publication/OTooleLRP21}, cites = {0}, citedby = {0}, pages = {26536-26547}, booktitle = {Advances in Neural Information Processing Systems 34: Annual Conference on Neural Information Processing Systems 2021, NeurIPS 2021, December 6-14, 2021, virtual}, editor = {Marc'Aurelio Ranzato and Alina Beygelzimer and Yann N. Dauphin and Percy Liang and Jennifer Wortman Vaughan}, }