{"id":{"repo_id":"wku-diss","oai_identifier":"oai:digitalcommons.wku.edu:theses-1062"},"canonical_url":"https://search.dev.ndltd.org/etd/wku-diss/oai:digitalcommons.wku.edu:theses-1062","repository":{"repo_id":"wku-diss","name":"Western Kentucky University","base_url":"https://digitalcommons.wku.edu/do/oai/"},"display":{"title":"A Framework for Consistency Based Feature Selection","abstract":"Feature selection is an effective technique in reducing the dimensionality of features in many applications where datasets involve hundreds or thousands of features. The objective of feature selection is to find an optimal subset of relevant features such that the feature size is reduced and understandability of a learning process is improved without significantly decreasing the overall accuracy and applicability. This thesis focuses on the consistency measure where a feature subset is consistent if there exists a set of instances of length more than two with the same feature values and the same class labels. This thesis introduces a new consistency-based algorithm, Automatic Hybrid Search (AHS) and reviews several existing feature selection algorithms (ES, PS and HS) which are based on the consistency rate. After that, we conclude this work by conducting an empirical study to a comparative analysis of different search algorithms.","abstract_html":"Feature selection is an effective technique in reducing the dimensionality of features in many applications where datasets involve hundreds or thousands of features. The objective of feature selection is to find an optimal subset of relevant features such that the feature size is reduced and understandability of a learning process is improved without significantly decreasing the overall accuracy and applicability. This thesis focuses on the consistency measure where a feature subset is consistent if there exists a set of instances of length more than two with the same feature values and the same class labels. This thesis introduces a new consistency-based algorithm, Automatic Hybrid Search (AHS) and reviews several existing feature selection algorithms (ES, PS and HS) which are based on the consistency rate. After that, we conclude this work by conducting an empirical study to a comparative analysis of different search algorithms.","abstract_has_math":false,"creators":["Lin, Pengpeng"],"institution":null,"degree_name":"Master of Science","degree_level":null,"degree_discipline":"Department of Mathematics and Computer Science","degree_department":null,"school":null,"contributors":["Dr. Huanjing Wang (Director),Dr. Art Shindhelm,Dr. Qi Li"],"advisors":[],"committee_chairs":[],"committee_members":[],"year":2009,"date_issued":"2009-05-01T07:00:00Z","date_published":"2009-05-01T07:00:00Z","updated_at":"2026-07-24T06:07:03Z","subjects":["data mining","probabilistic search algorithm","Automatic Hybrid search algorithm","Analysis","Mathematics"],"languages":[],"rights":[],"rights_urls":[],"identifier_entries":[]},"links":{"outbound_url":"https://digitalcommons.wku.edu/theses/62","outbound_label":"Repository record","outbound_source":"dc:identifier"},"metadata_groups":[{"id":"people","label":"People","entries":[{"key":"dc:contributor","label":"Contributor","values":["Dr. Huanjing Wang (Director),Dr. Art Shindhelm,Dr. Qi Li"]},{"key":"dc:creator","label":"Author","values":["Lin, Pengpeng"]}]},{"id":"academic_context","label":"Academic Context","entries":[{"key":"dc:type","label":"Dc Type","values":["Thesis"]},{"key":"thesis:degree_discipline","label":"Discipline","values":["Department of Mathematics and Computer Science"]},{"key":"thesis:degree_name","label":"Degree Name","values":["Master of Science"]}]},{"id":"subjects_keywords","label":"Subjects and Keywords","entries":[{"key":"dc:subject","label":"Dc Subject","values":["data mining","probabilistic search algorithm","Automatic Hybrid search algorithm","Analysis","Mathematics"]}]},{"id":"identifiers","label":"Identifiers","entries":[{"key":"dc:identifier","label":"Identifier","values":["https://digitalcommons.wku.edu/theses/62"]}]},{"id":"additional","label":"Additional Metadata","entries":[{"key":"dc:description.abstract","label":"Abstract","values":["Feature selection is an effective technique in reducing the dimensionality of features in many applications where datasets involve hundreds or thousands of features. The objective of feature selection is to find an optimal subset of relevant features such that the feature size is reduced and understandability of a learning process is improved without significantly decreasing the overall accuracy and applicability. This thesis focuses on the consistency measure where a feature subset is consistent if there exists a set of instances of length more than two with the same feature values and the same class labels. This thesis introduces a new consistency-based algorithm, Automatic Hybrid Search (AHS) and reviews several existing feature selection algorithms (ES, PS and HS) which are based on the consistency rate. After that, we conclude this work by conducting an empirical study to a comparative analysis of different search algorithms."]},{"key":"dc:title","label":"Title","values":["A Framework for Consistency Based Feature Selection"]}]}],"canonical_facts":{"dc:contributor":["Dr. Huanjing Wang (Director),Dr. Art Shindhelm,Dr. Qi Li"],"dc:creator":["Lin, Pengpeng"],"dc:description.abstract":["Feature selection is an effective technique in reducing the dimensionality of features in many applications where datasets involve hundreds or thousands of features. The objective of feature selection is to find an optimal subset of relevant features such that the feature size is reduced and understandability of a learning process is improved without significantly decreasing the overall accuracy and applicability. This thesis focuses on the consistency measure where a feature subset is consistent if there exists a set of instances of length more than two with the same feature values and the same class labels. This thesis introduces a new consistency-based algorithm, Automatic Hybrid Search (AHS) and reviews several existing feature selection algorithms (ES, PS and HS) which are based on the consistency rate. After that, we conclude this work by conducting an empirical study to a comparative analysis of different search algorithms."],"dc:identifier":["https://digitalcommons.wku.edu/theses/62"],"dc:subject":["data mining","probabilistic search algorithm","Automatic Hybrid search algorithm","Analysis","Mathematics"],"dc:title":["A Framework for Consistency Based Feature Selection"],"dc:type":["Thesis"],"thesis:degree_discipline":["Department of Mathematics and Computer Science"],"thesis:degree_name":["Master of Science"]},"updated_at":"2026-07-24T06:07:03Z"}