@inproceedings{24088abbb05a4fd6ae40455fc3a97902,
title = "Generating Synthetic Samples to Improve Small Sample Learning with Mixed Numerical and Categorical Attributes",
abstract = "The small data learning issue has existed for over one hundred years (since 1908) when the Student's t-distribution was first developed. Few statistical tools can evaluate a population appropriately if the sample size is too small; small samples can be remedied through virtual sample generation (VSG) methods, which are widely used in industry and machine learning. However, most VSG methods were developed for data having only numerical attributes, very few studies have dealt with nominal attributes and cause domain estimation limitations. Therefore, this paper proposes a method that generates virtual samples based on the discrete degrees of nominal attributes, and then estimates the general population domains by fuzzy membership functions. A backpropagation neural network model and a support vector regression model are used to test the efficiency of the proposed method, while the Wilcoxon-sign test is used to test the difference with raw data sets. The result shows that the proposed method can reduce the mean absolute error and enhance classification accuracy by generating virtual samples that have nominal attributes.",
author = "Lin, {Yao San} and Cheng, {Wan Ni} and Chen, {Chien Chih} and Li, {Der Chiang} and Chen, {Hung Yu}",
year = "2019",
month = jul,
doi = "10.1109/IIAI-AAI.2019.00121",
language = "English",
series = "Proceedings - 2019 8th International Congress on Advanced Applied Informatics, IIAI-AAI 2019",
publisher = "Institute of Electrical and Electronics Engineers Inc.",
pages = "567--572",
booktitle = "Proceedings - 2019 8th International Congress on Advanced Applied Informatics, IIAI-AAI 2019",
address = "United States",
note = "8th IIAI International Congress on Advanced Applied Informatics, IIAI-AAI 2019 ; Conference date: 07-07-2019 Through 11-07-2019",
}