@inproceedings{5405dc2800894c44a17ec59df8039865,
title = "SRes-NeRF: Improved Neural Radiance Fields for Realism and Accuracy of Specular Reflections",
abstract = "The Neural Radiance Fields (NeRF) is a popular view synthesis technique that represents a scene using a multilayer perceptron (MLP) combined with classic volume rendering and uses positional encoding techniques to increase image resolution. Although it can effectively represent the appearance of a scene, they often fail to accurately capture and reproduce the specular details of surfaces and require a lengthy training time ranging from hours to days for a single scene. We address this limitation by introducing a representation consisting of a density voxel grid and an enhanced MLP for a complex view-dependent appearance and model acceleration. Modeling with explicit and discretized volume representations is not new, but we propose Swish Residual MLP (SResMLP). Compared with the standard MLP+ReLU network, the introduction of layer scale module allows the shallow information of the network to be transmitted to the deep layer more accurately, maintaining the consistency of features. Introduce affine layers to stabilize training, accelerate convergence and use the Swish activation function instead of ReLU. Finally, an evaluation of four inward-facing benchmarks shows that our method surpasses NeRF{\textquoteright}s quality, it only takes about 18 min to train from scratch for a new scene and accuracy capture the specular details of the scene surface. Excellent performance even without positional encoding.",
keywords = "3D deep learning, Image-based rendering, Scene representation, Spectral bias, View synthesis, Volume rendering",
author = "Shufan Dai and Yangjie Cao and Pengsong Duan and Xianfu Chen",
note = "Publisher Copyright: {\textcopyright} 2023, The Author(s), under exclusive license to Springer Nature Switzerland AG.; 29th International Conference on MultiMedia Modeling, MMM 2023 ; Conference date: 09-01-2023 Through 12-01-2023",
year = "2023",
doi = "10.1007/978-3-031-27077-2_24",
language = "English",
isbn = "978-3-031-27076-5",
series = "Lecture Notes in Computer Science (including subseries Lecture Notes in Artificial Intelligence and Lecture Notes in Bioinformatics)",
publisher = "Springer",
pages = "306--317",
editor = "Duc-Tien Dang-Nguyen and Cathal Gurrin and Smeaton, {Alan F.} and Martha Larson and Stevan Rudinac and Minh-Son Dao and Christoph Trattner and Phoebe Chen",
booktitle = "MultiMedia Modeling: 29th International Conference, MMM 2023",
address = "Germany",
}