@inproceedings{61aa2567de934e1c8831217197c2f10f,
title = "Using machine learning to predict ranking of webpages in the gift industry: Factors for search-engine optimization",
abstract = "We use machine learning to predict the search engine rank of webpages. We use a list of keywords for 30 content blogs of an e-commerce company in the gift industry to retrieve 733 content pages occupying the first-page Google rankings and predict their rank using 30 ranking factors. We test two models, Light Gradient Boosting Machine (LightGBM) and Extreme Gradient Boosted Decision Trees (XGBoost), finding that XGBoost performs better for predicting actual search rankings, with an average accuracy of 0.86. The feature analysis shows the most impactful features are (a) internal and external links, (b) security of the web domain, and (c) length of H3 headings, and the least impactful features are (a) keyword mentioned in domain address, (b) keyword mentioned in the H1 headings, and (c) overall number of keyword mentions in the text. The results highlight the persistent importance of links in search-engine optimization. We provide actionable insights for online marketers and content creators.",
keywords = "Content Marketing, E-Commerce, Machine Learning, Online Marketing, Rank Prediction, Search-Engine Optimization",
author = "Joni Salminen and Roope Marttila and Jansen, {Bernard J.} and Juan Corporan and Tommi Salenius",
note = "Publisher Copyright: {\textcopyright} 2019 Copyright is held by the owner/author(s). Publication rights licensed to ACM.; 9th International Conference on Information Systems and Technologies, ICIST 2019 ; Conference date: 24-03-2019 Through 26-03-2019",
year = "2019",
month = mar,
day = "24",
doi = "10.1145/3361570.3361578",
language = "English",
series = "ACM International Conference Proceeding Series",
publisher = "Association for Computing Machinery",
booktitle = "Proceedings of the 9th International Conference on Information Systems and Technologies, ICIST 2019",
address = "United States",
}