<!DOCTYPE article PUBLIC "-//NLM//DTD JATS (Z39.96) Journal Archiving and Interchange DTD v1.0 20120330//EN" "JATS-archivearticle1.dtd">
<article xmlns:xlink="http://www.w3.org/1999/xlink">
  <front>
    <journal-meta>
      <journal-title-group>
        <journal-title>ITADATA</journal-title>
      </journal-title-group>
    </journal-meta>
    <article-meta>
      <contrib-group>
        <contrib contrib-type="author">
          <string-name>Bahne Christiansen</string-name>
          <email>bahne.christiansen@nordakademie.de</email>
          <xref ref-type="aff" rid="aff5">5</xref>
        </contrib>
        <contrib contrib-type="author">
          <string-name>Uwe Neuhaus</string-name>
          <email>uwe.neuhaus@uni-flensburg.de</email>
          <xref ref-type="aff" rid="aff2">2</xref>
        </contrib>
        <contrib contrib-type="author">
          <string-name>Michael Schulz</string-name>
          <email>michael.schulz@nordakademie.de</email>
          <xref ref-type="aff" rid="aff5">5</xref>
        </contrib>
        <contrib contrib-type="author">
          <string-name>Adrian Hargreaves</string-name>
        </contrib>
        <contrib contrib-type="author">
          <string-name>Antinisca Di</string-name>
        </contrib>
        <contrib contrib-type="author">
          <string-name>Marco</string-name>
          <xref ref-type="aff" rid="aff1">1</xref>
        </contrib>
        <contrib contrib-type="author">
          <string-name>Guido Proietti</string-name>
          <email>guido.proietti@univaq.it</email>
          <xref ref-type="aff" rid="aff1">1</xref>
        </contrib>
        <contrib contrib-type="author">
          <string-name>Fabrizio Rossi</string-name>
          <email>fabrizio.rossi@univaq.it</email>
          <xref ref-type="aff" rid="aff1">1</xref>
        </contrib>
        <contrib contrib-type="author">
          <string-name>Giovanni Stilo</string-name>
          <email>giovanni.stilo@univaq.it</email>
          <xref ref-type="aff" rid="aff1">1</xref>
        </contrib>
        <contrib contrib-type="author">
          <string-name>Mark H. Haney</string-name>
          <email>m.haney@chatham.edu</email>
          <xref ref-type="aff" rid="aff0">0</xref>
        </contrib>
        <contrib contrib-type="author">
          <string-name>Andrew Duncan</string-name>
        </contrib>
        <contrib contrib-type="author">
          <string-name>Daniele Tessera</string-name>
          <email>daniele.tessera@unicatt.it</email>
          <xref ref-type="aff" rid="aff6">6</xref>
        </contrib>
        <contrib contrib-type="author">
          <string-name>Tilman Todt</string-name>
          <email>tilman.todt@han.nl</email>
          <xref ref-type="aff" rid="aff3">3</xref>
        </contrib>
        <contrib contrib-type="author">
          <string-name>Andreas Brandenberg</string-name>
          <email>andreas.brandenberg@hslu.ch</email>
          <xref ref-type="aff" rid="aff4">4</xref>
        </contrib>
        <contrib contrib-type="author">
          <string-name>Patricia Feubli</string-name>
          <email>patricia.feubli@hslu.ch</email>
          <xref ref-type="aff" rid="aff4">4</xref>
        </contrib>
        <contrib contrib-type="author">
          <string-name>Whitireia New Zealand</string-name>
        </contrib>
        <contrib contrib-type="author">
          <string-name>Wellington</string-name>
        </contrib>
        <contrib contrib-type="author">
          <string-name>New Zealand</string-name>
        </contrib>
        <contrib contrib-type="author">
          <string-name>Inverness College UHI</string-name>
        </contrib>
        <contrib contrib-type="author">
          <string-name>Inverness</string-name>
        </contrib>
        <contrib contrib-type="author">
          <string-name>Scotland</string-name>
        </contrib>
        <aff id="aff0">
          <label>0</label>
          <institution>Chatham University</institution>
          ,
          <addr-line>Pittsburgh</addr-line>
          ,
          <country country="US">United States of America</country>
        </aff>
        <aff id="aff1">
          <label>1</label>
          <institution>Dipartimento di Ingegneria e Scienze dell'Informazione e Matematica University of L'Aquila</institution>
          ,
          <addr-line>L'Aquila</addr-line>
          ,
          <country country="IT">Italy</country>
        </aff>
        <aff id="aff2">
          <label>2</label>
          <institution>Europa-Universität Flensburg</institution>
          ,
          <addr-line>Flensburg</addr-line>
          ,
          <country country="DE">Germany</country>
        </aff>
        <aff id="aff3">
          <label>3</label>
          <institution>Hogeschool van Arnhem en Nijmegen</institution>
          ,
          <addr-line>Arnhem</addr-line>
          ,
          <country country="NL">Netherlands</country>
        </aff>
        <aff id="aff4">
          <label>4</label>
          <institution>Lucerne University of Applied Sciences and Arts</institution>
          ,
          <addr-line>Lucerne</addr-line>
          ,
          <country country="CH">Switzerland</country>
        </aff>
        <aff id="aff5">
          <label>5</label>
          <institution>NORDAKADEMIE Hochschule der Wirtschaft</institution>
          ,
          <addr-line>Elmshorn</addr-line>
          ,
          <country country="DE">Germany</country>
        </aff>
        <aff id="aff6">
          <label>6</label>
          <institution>Università Cattolica del Sacro Cuore</institution>
          ,
          <addr-line>Brescia</addr-line>
          ,
          <country country="IT">Italy</country>
        </aff>
      </contrib-group>
      <pub-date>
        <year>2022</year>
      </pub-date>
      <volume>1</volume>
      <fpage>20</fpage>
      <lpage>21</lpage>
      <abstract>
        <p>Due to the increasing growth of the discipline Data Science, a partition into sub-disciplines seems appropriate. Therefore, we propose the division of Data Science in Pure Data Science, in which methods and tools are developed, and Applied Data Science, in which these methods and tools are adapted and applied to practical problems of a specific domain. This article focuses on Applied Data Science and how it should be positioned in relation to its adjoining disciplines. We also introduce the term Business Data Science as a specific form of Applied Data Science in the business domain and describe its relationship to existing terms like</p>
      </abstract>
      <kwd-group>
        <kwd>Data Science</kwd>
        <kwd>Applied Data Science</kwd>
        <kwd>Pure Data Science</kwd>
        <kwd>Business Data Science</kwd>
      </kwd-group>
    </article-meta>
  </front>
  <body>
    <sec id="sec-1">
      <title>1. Introduction</title>
      <p>
        Data Science (DS) is developing into an independent new subject area. Highly interdisciplinary by
nature, it takes concepts and methods from mathematics/statistics and computer science, it combines,
expands, and enhances them and applies them to new areas of application [
        <xref ref-type="bibr" rid="ref1">1</xref>
        ]. As an emerging new
subject area, DS also develops its own research questions, processes, and techniques, independent of
its underlying disciplines [
        <xref ref-type="bibr" rid="ref2">2</xref>
        ].
      </p>
      <p>
        Although DS now possesses generic methods and algorithms that can be applied in many domains,
they often must be refined to the particular requirements of domain-specific data applications [
        <xref ref-type="bibr" rid="ref3">3</xref>
        ].
Systematic approaches are needed to address the complexities of DS problems inherent within these
domains, which are not effectively accommodated within a single discipline [
        <xref ref-type="bibr" rid="ref4">4</xref>
        ].
      </p>
      <p>2022 Copyright for this paper by its authors.</p>
      <p>To be able to address the aforementioned aspects more concretely, the focus of this article will be
on the application of DS to the business domain. This seems appropriate since DS is already widely
used in this domain. The authors work for different universities that have independently of each other
developed teaching programs that emphasize the application of DS. Although we had never heard of
each other before, we have discovered many commonalities in these programs. Through structured
expert discussions over six iterations, we extracted this common view. The results given here represent
a proposal for structuring the DS discipline.</p>
      <p>As a first result, we suggest dividing DS into two sub-disciplines; a sub-discipline in which methods
and tools are developed, and another sub-discipline in which these methods and tools are adapted and
applied to practical problems of a specific domain. Following the conventions in Mathematics, where
a similar distinction is made, we propose the terms Pure Data Science (PDS) to identify the
subdisciplines focused on the further development of the scientific field of the DS and Applied Data
Science (ADS) which is mainly devoted to solving application domain challenges and provide an
explanation of the obtained results. This manifesto focuses on ADS and how it should be positioned in
relation to its adjoining disciplines.</p>
      <p>An overview of the sub-disciplines (i.e. PDS versus ADS) with their distinguished characteristics is
shown in</p>
      <p>Table 1. The distinguishing characteristics are the main goal of the sub-discipline, which identifies
its purpose, the expected output (i.e. its expected outcome) of it, its primary beneficiary area, and the
required competencies needed to develop it further.</p>
      <p>A closer look at the degree programs named Applied Data Science shows that even these have
distinctly different emphases, so that a further subdivision of the subdiscipline should be made, which
will be presented in the following chapters. In this manifesto, we focus on the application of DS to
business problems. There are two main reasons for this: First, a discussion of ADS is only partially
possible without considering an explicit application domain. Second, the chosen application domain is
particularly relevant due to its widespread use. However, it is also an area that requires differentiation
from existing terms, such as business intelligence and business analytics, which are also described in
this manifesto.</p>
      <p>The paper is organized as follows. In Section 2 we present the methodology that shaped this
manifesto. In Section 3 we discuss Applied Data Science in the context of its adjoining relevant
disciplines by showing their respective influences. In Section 4 we focus our discussion on Business
Data Science, while in Section 5 we close with a summary of our findings.</p>
    </sec>
    <sec id="sec-2">
      <title>2. Methodology</title>
      <p>Our research shows that since 2020, more and more ADS degree programs are being founded
worldwide. The simultaneous birth of many identically named study programs from different
universities can hardly be a coincidence. A search in early November 2021 identified 60 of such
programs. Thus, we believe that DS is an emerging discipline with multiple facets due to its
multidisciplinary nature and an increasing number of application domains. The adoption of DS to
application domains shapes DS into sub-disciplines having a common core but requiring different
knowledge and skills tailored to the domains where it is applied. Moreover, such application induces
cross-fertilization between DS and the considered application domain shaping novel competencies and,
in some cases novel disciplines.</p>
      <p>In view of the above considerations, a group of international academics and lecturers, which offer
degree programs in ADS at different universities, began a discussion on emerging DS-related
disciplines and started to define a taxonomy to be used to classify research directions as well as teaching
projects, and to shape the emerging DS sub-disciplines.</p>
      <p>The group worked iteratively on this manifesto, starting from the individual degree programs and
their experiences. At each iteration, they focused on a specific aspect and worked on it until reaching
an agreement among the authors. Finally, they concentrated on the business domain since all authors
have a common interest in it.</p>
    </sec>
    <sec id="sec-3">
      <title>3. Applied Data Science</title>
      <p>Having discussed the division of Data Science into the sub-disciplines of PDS and ADS, we now
move ADS to the center of our discussion.</p>
      <p>In particular, we consider ADS in the context of its adjoining relevant disciplines by showing their
respective influences. Figure 1 provides an overview of the relations of ADS with: (1) PDS and its
related disciplines (2) Focus dependent disciplines, (3) Domain and (4) Ethics and Regulations.</p>
      <sec id="sec-3-1">
        <title>Applied</title>
      </sec>
      <sec id="sec-3-2">
        <title>Data Science</title>
        <p>The main purpose of ADS is to solve real-world problems in a domain using and possibly adapting
methods and tools provided by PDS. Depending on the focus of an individual project, other disciplines
such as linguistics or psychology may also be relevant.</p>
        <p>The connections between Data Science and neighboring disciplines are often represented as a Venn
diagram with Data Science as the intersection of other disciplines. This type of visualization incorrectly
suggests that Data Science is merely a compilation of existing concepts from surrounding disciplines
and does not adequately represent the increasing number of concepts and methods developed within
Data Science – or with a more detailed view, within PDS and ADS. The type of diagram chosen in our
figure is intended to highlight precisely this fact, which justifies the definition of a (sub)discipline in
the first place. Self-developed methods of ADS include re-usable domain-specific strategies for
applying data mining techniques, communication processes with decision makers, managing analytical
projects, or quality attributes as privacy, explainability, debias, fairness, replicability, accuracy, etc.
These aspects could be added to the center of the above figure. However, since a complete enumeration
is not possible, we have refrained from doing so.</p>
        <p>The domain-independent methods and tools developed by PDS rely heavily on the disciplines of
mathematics/statistics, computer science and information science. Depending on the specific
application domain, other specific disciplines will provide additional support. Therefore, all these
disciplines play an important role in ADS as well.</p>
        <p>We illustrate the relationships of the four neighboring areas to ADS in an example. Consider an
insurance company that wants to identify fraudulent claims. For this purpose, knowledge from all four
neighboring areas of ADS are relevant: From the domain (insurance) the contract-specific insurance
terms and conditions, from the focus dependent disciplines (in this example linguistics) relevant
methods for text analysis, from Ethics and Regulations the consideration, if social customer features
should be included, and from Pure Data Science a suitable classification algorithm. Furthermore,
ADSspecific concepts such as explainability and debiasing are relevant.</p>
        <sec id="sec-3-2-1">
          <title>Domain-specific ADS forms</title>
          <p>Having clarified that ADS is a sub-discipline of DS, a natural question is whether ADS also has
subdisciplines itself. The formation of an independent sub-discipline of ADS seems appropriate to us if the
following criteria are met:
(a) The application domain is sufficiently large.
(b) There are many questions in the application domain that can be answered with the help of ADS.
(c) The questions to be answered make special adjustments and further developments of the ADS
methods necessary.</p>
          <p>Criteria (a) and (b) are intended to ensure that the application domain under consideration is
significant enough to justify the creation of an independent ADS sub-discipline (with a specialized
training program, special publications, etc.). Criterion (c) is to ensure that the domain questions cannot
be answered using standard ADS methods alone. If that were the case, the creation of a dedicated
subdiscipline would not be necessary.</p>
          <p>There are a variety of application domains that, according to the above criteria, justify the creation
of a separate sub-discipline. Besides the domain they differ in the following aspects:
• focus on specific types of data (structured databases, text data, sensor data, audio/video data)
• use different subsets of the “PDS toolbox”
• adapt general PDS methods to the needs of the domain
• may develop domain specific DS methods
• may use specialized tools/software that supports domain specific DS methods
• take into account domain specific requirements, e.g.</p>
          <p>o legal and ethical concerns
o data privacy
o organizational standards
o scientific rigor
o explainability of results
o fairness/bias
o performance</p>
          <p>While a complete listing of all domain-specific ADS forms is difficult, a subdivision into the two
forms ADS in a scientific domain (e.g., Biological Data Science, Climate Data Science, and in more
general sense, Experimental Sciences) and ADS in a non-scientific domain (e.g., Business Data Science,
Public Service Data Science, Engineering Data Science) seems useful due to the differences presented
below. Table 2 shows some of the differences we consider between ADS in a scientific domain and
ADS in a non-scientific domain.</p>
        </sec>
      </sec>
    </sec>
    <sec id="sec-4">
      <title>4. Business Data Science</title>
      <p>Applied Data Science in a
scientific domain
Scientific insights
Often low
Often deep
Typically conducted by the
domain experts (scientists)
themselves</p>
      <p>Applied Data Science in a
nonscientific domain
Pragmatic insights and their
usage
Often high
Mostly medium
Typically conducted by Applied
Data Scientists</p>
      <p>In contrast to the previous considerations, the classification of terms in the context of analytics with
a business perspective has already been considered intensively in the literature, but a consensus has not
yet been found. We first give an overview of the common usages of the terms Business Analytics (BA)
and Business Intelligence (BI). Then we introduce the term Business Data Science and describe its
relationship to BA/BI.</p>
      <p>
        The literature reveals significant areas of ambiguity in relation to the demarcations that distinguish
the disciplines and subfields that concern organizational decision making. For example, some authors
still do not distinguish between the terms Business Analytics and Data Science [
        <xref ref-type="bibr" rid="ref5 ref6 ref7 ref8">5, 6, 7, 8</xref>
        ]. However,
many other authors do distinguish the two [
        <xref ref-type="bibr" rid="ref10 ref11 ref9">9, 10, 11</xref>
        ]. Despite this ambiguity, reviews of the literature
have identified emerging themes relating to the demarcation of relevant disciplines.
      </p>
      <p>
        A characteristic that is often used in the literature to delimit the terms is the type of statistical
methods used. It is therefore advantageous to first break down the statistics into its three sub-disciplines
[
        <xref ref-type="bibr" rid="ref12">12</xref>
        ]:
• Descriptive statistics is used to describe a data set (representing a phenomena) using
meaningful parameters such as location, dispersion, or correlation measures as well as the graphical
visualization of the data. Such techniques are used to analyze operational business data to aggregate
decision-supporting information by means of data-based dashboards, scorecards, or reports. In
descriptive statistics, neither hypothesis tests based on mathematical models are used to generalize
the results nor are probabilistic forecasts made. Descriptive statistics does not use probability theory.
• Explorative statistics deals with the extraction of structures from data. It uses the results of
descriptive statistics and creates mathematical models or statistical hypotheses. Many analytical
methods of explorative statistics (e.g. cluster analyses, factor analyses, principal component
analyses) are used in data mining to reduce the complexity of large amounts of data and thus gain
insight into underlying relationships [
        <xref ref-type="bibr" rid="ref13">13</xref>
        ].
• Inferential statistics tests the models and hypotheses generated from data by exploratory
algorithms and makes intensive use of the probability theory. The aims are to produce reliable
generalizations of results beyond the existing data set or forecast future developments. The
application of such methods in a business context is usually summarized as predictive analytics.
      </p>
      <p>
        Analytics has its origins within logic, mathematics, and science. Grammatically, the term Analytics
includes the suffix “-ics”, which refers to a body of knowledge or principles. Nelson defines Analytics
as “the scientific process or discipline of fact-based problem-solving” [
        <xref ref-type="bibr" rid="ref14">14</xref>
        ]. Davenport and Harris define
Analytics as the “extensive use of data, statistical and quantitative analysis, exploratory and predictive
models, and fact-based management to drive decisions and actions” [
        <xref ref-type="bibr" rid="ref15">15</xref>
        ]. There are many other
definitions within the literature, but this theme of applying advanced analytics and statistical techniques
on data to drive decision making is often a common thread that links them together.
      </p>
      <p>
        Analytics is a broad discipline that logically includes the sub-fields of BA and Data Science.
Business Analytics is primarily concerned with business relevance and actionable insights from
analyses. This concern is highlighted in a review performed by Phelps and Szabat [
        <xref ref-type="bibr" rid="ref9">9</xref>
        ]. They found that
definitions of BA often contained a substantial focus on the statistical and quantitative analysis of data,
and decision-making support within business domains. Their review also sought definitions of Data
Science from the literature and they observed clear differences between the two. The definitions they
found relating to Data Science include four key aspects: data (data modeling, taxonomy, data
management, data optimization, onthology, ethical and legal usage of data, etc.), databases, computer
systems (transformation of inputs into outputs) and advanced analytics including statistics. The degree
of advanced analytics that each of these typically employs is, at least, one significant difference between
these two sub-fields. The utilization of advanced analytics in Data Science additionally demands
methods for quality assurance or improvement (privacy, debiasing, fairness, accuracy, etc.) and as a
result more advanced tools than in BA. This is also true for statistics. While BA does apply explorative
statistics frequently, the use of inferential statistics is more common in Data Science.
      </p>
      <p>
        Previously, and to some extent currently, BI has been described in the literature as an umbrella term
that includes the full range of available strategies and technologies that enable data driven decision
making (e.g., [
        <xref ref-type="bibr" rid="ref16">16</xref>
        ]). Based on this very broad definition, BA could be considered as subfields of BI.
However, other authors (e.g., [
        <xref ref-type="bibr" rid="ref17">17</xref>
        ]) define BI more narrowly, as technologies that apply descriptive
statistics to improve strategic decision making. Certainly, within industry, BI tools used to implement
what are referred to as Business Intelligence solutions provide functionality that focuses primarily on
summarizing and visualizing the data found in data warehouses, using technologies such as OLAP
(online analytical processing). Typically, these tools analyze data using descriptive statistical
techniques and provide decision makers with visual tools to monitor business performance against a
variety of KPIs. Using this narrower definition of BI, it can be considered as a sub-field of BA due to
its application of descriptive statistics to aid decision making.
      </p>
      <p>
        Taking a business perspective, the term Big Data Analytics has been described as a discipline that
has emerged from BI [
        <xref ref-type="bibr" rid="ref18">18</xref>
        ]. Big Data Analytics concerns the organizing of big data, analyzing and
discovering knowledge, patterns and intelligence from big data, visualization and the reporting of
discovered knowledge for assisting decision making [
        <xref ref-type="bibr" rid="ref19">19</xref>
        ]. The authors claim that the main components
of big data analytics include descriptive statistics, predictive analytics and prescriptive analytics.
      </p>
      <sec id="sec-4-1">
        <title>Introducing the term Business Data Science</title>
        <p>
          Analogous to the term Business Informatics in Computer Science, we propose to introduce the term
Business Data Science (BDS). In our opinion, BDS is a sub-discipline of ADS, in which the methods
and instruments of ADS are applied to issues in the business domain. Methods and instruments are
adapted to the requirements of the domain and, if necessary, expanded. BDS can therefore be defined
without reference to BI and BA. It is sufficient to concretize a general data science definition (e.g. [
          <xref ref-type="bibr" rid="ref20">20</xref>
          ])
for application to the business domain:
        </p>
        <p>Business Data Science is a field of interdisciplinary expertise in which scientific procedures are
used to (semi)automatically generate business insights from conceivably complex data leveraging
existing or newly developed analysis methods. The gained business insights are subsequently utilized
mainly for decision support, taking into account the effects on society.</p>
        <p>
          The business domain is one of the most important application areas of DS [
          <xref ref-type="bibr" rid="ref21">21</xref>
          ], and ADS addresses
a wide variety of business-related problems, many of which require modified or even specifically
created ADS methods and tools. The establishment of a separate sub-discipline of ADS therefore seems
justified. Though the term Business Data Science is not used frequently yet, it appears in a growing
number of publications (e.g., [
          <xref ref-type="bibr" rid="ref22">22, 23, 24</xref>
          ].
        </p>
        <p>The relationship between BI, BA and BDS can be represented as follows. As described above, the
aim of BI and BA is to support data-driven decisions in the business domain. In contrast to BI, BA also
uses advanced and complex methods of statistics, information science and computer science.
Historically, BI is the older discipline that was later supplemented/expanded by BA. BA continues to
grow strongly as more and more questions in the business domain are being addressed with increasingly
complex methods.</p>
        <p>The goals and methods of BA and BDS are very similar and differ more in their point of view than
in their content. BA represents more the internal perspective of the business domain. It focuses on the
business benefits and sees the DS primarily as an auxiliary science. BDS, on the other hand, represents
more of a domain-external perspective, focuses on the correct application of methods and sees the DS
as a guiding discipline. BA and BDS tend to develop towards each other, so that the terms can be used
synonymously in perspective.</p>
      </sec>
    </sec>
    <sec id="sec-5">
      <title>5. Summary</title>
      <p>Data Science can be divided into two sub-disciplines: Pure Date Science, in which tools and methods
are developed, and Applied Data Science, in which these tools and methods are applied and adapted to
practical problems of a specific domain. The focus of this manifesto is on the latter.</p>
      <p>We have presented ADS in the context of its four neighboring disciplines: (1) PDS, (2) Focus
dependent disciplines, (3) Domain and (4) Ethics &amp; Regulations. It should be emphasized that ADS is
more than the intersection of the neighboring disciplines. This aspect was illustrated by an appropriate
graphical visualization in Figure 1 and by a concrete example.</p>
      <p>Analogous to the subdivision of DS into the subdisciplines PDS and ADS, ADS can be decomposed
into numerous domain-specific subdisciplines such as Business Data Science. It is useful – regarding
their separating characteristics - to classify these sub-disciplines into two categories: ADS in a scientific
domain and ADS in a non-scientific domain by considering the specific challenges ADS aims to solve.</p>
      <p>An important sub-discipline of ADS focusing on a non-scientific domain is BDS. Traditionally, the
terms BI and BA are often used for data-driven analytical activities in the business context. While BDS
can be easily distinguished from BI, there are major methodological similarities to BA. The main
difference is the perspective on the problem to be analyzed. Besides BDS, other sub-disciplines of ADS
exist such as Public Service Data Science or Engineering Data Science. A closer look at these
subdisciplines was not the focus of this manifesto.</p>
      <p>The relationships between the various sub-disciplines have been summarized in Figure 2.</p>
      <p>Data Science
Pure Data Science</p>
      <p>Applied Data Science
Applied Data Science in a
scientific domain</p>
      <p>Applied Data Science in a</p>
      <p>non-scientific domain
Business Data Science
…</p>
      <p>In the future, to improve the presented classification and to enhance the developed taxonomy in
order to better define emerging DS sub-disciplines with all their facets, we plan to extend the survey to
a broad group of researchers and practitioners working in all areas of data science.</p>
      <p>As a final remark, we believe DS will become even more pervasive and will have applications that
we still do not imagine.</p>
    </sec>
    <sec id="sec-6">
      <title>6. References</title>
      <p>[23] Davenport, T., “Beyond unicorns: Educating, classifying, and certifying business data scientists”,</p>
      <p>Harvard Data Science Review, 2020.
[24] Miah, S. J., Solomonides, I., and Gammack, J. G., “A design-based research approach for
developing data-focussed business curricula”, Education and Information Technologies 25.1
(2020): 553–581.</p>
    </sec>
  </body>
  <back>
    <ref-list>
      <ref id="ref1">
        <mixed-citation>
          [1]
          <string-name>
            <surname>Blei</surname>
            ,
            <given-names>D. M.</given-names>
          </string-name>
          , and
          <string-name>
            <surname>Smyth</surname>
            ,
            <given-names>P.</given-names>
          </string-name>
          , “
          <article-title>Science and data science”</article-title>
          ,
          <source>Proceedings of the National Academy of Sciences</source>
          ,
          <volume>114</volume>
          .33 (
          <year>2017</year>
          ):
          <fpage>8689</fpage>
          -
          <lpage>8692</lpage>
          .
        </mixed-citation>
      </ref>
      <ref id="ref2">
        <mixed-citation>
          [2]
          <string-name>
            <surname>Braschler</surname>
            ,
            <given-names>M.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Stadelmann</surname>
            ,
            <given-names>T.</given-names>
          </string-name>
          , and
          <string-name>
            <surname>Stockinger</surname>
            ,
            <given-names>K.</given-names>
          </string-name>
          , “Data science”, in: Braschler M.,
          <string-name>
            <surname>Stadelmann</surname>
            <given-names>T.</given-names>
          </string-name>
          and
          <string-name>
            <surname>Stockinger</surname>
            <given-names>K</given-names>
          </string-name>
          . (eds) Applied Data Science. Springer, Cham,
          <year>2019</year>
          .
        </mixed-citation>
      </ref>
      <ref id="ref3">
        <mixed-citation>
          [3]
          <string-name>
            <surname>Brodie</surname>
            ,
            <given-names>M.</given-names>
          </string-name>
          , “What Is Data Science?”, In: Braschler M.,
          <string-name>
            <surname>Stadelmann</surname>
            <given-names>T.</given-names>
          </string-name>
          and
          <string-name>
            <surname>Stockinger</surname>
            <given-names>K</given-names>
          </string-name>
          . (eds) Applied Data Science. Springer, Cham,
          <year>2019</year>
          .
        </mixed-citation>
      </ref>
      <ref id="ref4">
        <mixed-citation>
          [4]
          <string-name>
            <surname>Longbing</surname>
            ,
            <given-names>C.</given-names>
          </string-name>
          ,
          <article-title>“Data Science: A Comprehensive Overview”</article-title>
          ,
          <source>ACM Computing Surveys</source>
          .
          <volume>50</volume>
          (
          <year>2017</year>
          ):
          <fpage>1</fpage>
          -
          <lpage>42</lpage>
          .
        </mixed-citation>
      </ref>
      <ref id="ref5">
        <mixed-citation>
          [5]
          <string-name>
            <surname>Henry</surname>
            ,
            <given-names>R.</given-names>
          </string-name>
          and
          <string-name>
            <surname>Venkatraman</surname>
            ,
            <given-names>S.</given-names>
          </string-name>
          , “
          <article-title>Big Data Analytics the Next Big Learning Opportunity”</article-title>
          ,
          <source>Journal of Management Information and Decision Sciences</source>
          ,
          <volume>18</volume>
          .2 (
          <year>2015</year>
          ):
          <fpage>17</fpage>
          -
          <lpage>29</lpage>
          .
        </mixed-citation>
      </ref>
      <ref id="ref6">
        <mixed-citation>
          [6]
          <string-name>
            <surname>Miller</surname>
            ,
            <given-names>S.</given-names>
          </string-name>
          , “
          <article-title>Collaborative Approaches Needed to Close the Big Data Skills Gap”</article-title>
          ,
          <source>Journal of Organization Design</source>
          ,
          <volume>3</volume>
          .1 (
          <year>2014</year>
          ):
          <fpage>26</fpage>
          -
          <lpage>30</lpage>
          .
        </mixed-citation>
      </ref>
      <ref id="ref7">
        <mixed-citation>
          [7]
          <string-name>
            <surname>Power</surname>
            ,
            <given-names>D.J.</given-names>
          </string-name>
          , “
          <article-title>Data science: supporting decision-making”</article-title>
          ,
          <source>Journal of Decision Systems</source>
          ,
          <volume>25</volume>
          .4 (
          <year>2016</year>
          ):
          <fpage>345</fpage>
          -
          <lpage>356</lpage>
          .
        </mixed-citation>
      </ref>
      <ref id="ref8">
        <mixed-citation>
          [8]
          <string-name>
            <surname>Kemper</surname>
            ,
            <given-names>S.</given-names>
          </string-name>
          and
          <string-name>
            <surname>Mathews</surname>
            ,
            <given-names>T.</given-names>
          </string-name>
          ,
          <source>“Earth Science Data Analytics: Definitions, Techniques and Skills”, Data Science Journal</source>
          ,
          <volume>16</volume>
          .6 (
          <year>2017</year>
          ):
          <fpage>1</fpage>
          -
          <lpage>8</lpage>
          .
        </mixed-citation>
      </ref>
      <ref id="ref9">
        <mixed-citation>
          [9]
          <string-name>
            <surname>Phelps</surname>
            ,
            <given-names>A.L.</given-names>
          </string-name>
          and
          <string-name>
            <surname>Szabat</surname>
            ,
            <given-names>K.A.</given-names>
          </string-name>
          , “
          <article-title>The Current Landscape of Teaching Analytics to Business Students at Institutions of Higher Education: Who is Teaching What?”</article-title>
          , The American Statistician,
          <volume>71</volume>
          .2 (
          <year>2017</year>
          ).
        </mixed-citation>
      </ref>
      <ref id="ref10">
        <mixed-citation>
          [10]
          <string-name>
            <surname>Bichler</surname>
            ,
            <given-names>M.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Heinzl</surname>
            ,
            <given-names>A.</given-names>
          </string-name>
          , and
          <string-name>
            <surname>Van der Aalst</surname>
          </string-name>
          , W.M., “Business Analytics and Data Science: Once Again?”,
          <source>Business &amp; Information Systems Engineering</source>
          ,
          <volume>59</volume>
          .2 (
          <year>2017</year>
          ):
          <fpage>77</fpage>
          -
          <lpage>79</lpage>
          .
        </mixed-citation>
      </ref>
      <ref id="ref11">
        <mixed-citation>
          [11]
          <string-name>
            <surname>Kambatla</surname>
            ,
            <given-names>K.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Kollias</surname>
            ,
            <given-names>G.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Kumar</surname>
            ,
            <given-names>V.</given-names>
          </string-name>
          and
          <string-name>
            <surname>Grama</surname>
            ,
            <given-names>A.</given-names>
          </string-name>
          , “
          <article-title>Trends in big data analytics”</article-title>
          ,
          <source>Journal of Parallel and Distributed Computing</source>
          ,
          <volume>74</volume>
          (
          <year>2014</year>
          ):
          <fpage>2561</fpage>
          -
          <lpage>2573</lpage>
          .
        </mixed-citation>
      </ref>
      <ref id="ref12">
        <mixed-citation>
          [12]
          <string-name>
            <surname>Laursen</surname>
            ,
            <given-names>G. H.</given-names>
          </string-name>
          , and
          <string-name>
            <surname>Thorlund</surname>
          </string-name>
          , J., “
          <article-title>Business analytics for managers: Taking business intelligence beyond reporting”</article-title>
          , John Wiley &amp; Sons,
          <year>2016</year>
          .
        </mixed-citation>
      </ref>
      <ref id="ref13">
        <mixed-citation>
          [13]
          <string-name>
            <surname>Wixom</surname>
            ,
            <given-names>B.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Ariyachandra</surname>
            ,
            <given-names>T.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Douglas</surname>
            ,
            <given-names>D.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Goul</surname>
            ,
            <given-names>M.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Gupta</surname>
            ,
            <given-names>B.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Iyer</surname>
            ,
            <given-names>L.</given-names>
          </string-name>
          , and
          <string-name>
            <surname>Turetken</surname>
            ,
            <given-names>O.</given-names>
          </string-name>
          , “
          <article-title>The current state of business intelligence in academia: The arrival of big data”</article-title>
          ,
          <source>Communications of the Association for information Systems</source>
          ,
          <volume>34</volume>
          .1 (
          <year>2014</year>
          ).
        </mixed-citation>
      </ref>
      <ref id="ref14">
        <mixed-citation>
          [14]
          <string-name>
            <surname>Nelson</surname>
            ,
            <given-names>G.</given-names>
          </string-name>
          ,
          <article-title>Difference between analytics and big data, data science and informatics</article-title>
          .
          <source>ThotWave Blog</source>
          ,
          <year>2017</year>
          . Retrieved from https://www.thotwave.com/blog/2017/07/07/difference-betweenanalytics-and
          <string-name>
            <surname>-</surname>
          </string-name>
          bigdatadatascience-informatics/
        </mixed-citation>
      </ref>
      <ref id="ref15">
        <mixed-citation>
          [15]
          <string-name>
            <surname>Davenport</surname>
            ,
            <given-names>T. H.</given-names>
          </string-name>
          , and
          <string-name>
            <surname>Harris</surname>
            ,
            <given-names>J. G.</given-names>
          </string-name>
          ,
          <article-title>“Competing on analytics: The new science of winning”</article-title>
          , MA: Harvard Business School Press,
          <year>2007</year>
          .
        </mixed-citation>
      </ref>
      <ref id="ref16">
        <mixed-citation>
          [16]
          <string-name>
            <surname>Turban</surname>
            ,
            <given-names>E.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Sharda</surname>
            ,
            <given-names>R.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Delen</surname>
            ,
            <given-names>D.</given-names>
          </string-name>
          and
          <string-name>
            <surname>King</surname>
            ,
            <given-names>D.</given-names>
          </string-name>
          , “Business Intelligence:
          <string-name>
            <given-names>A Managerial</given-names>
            <surname>Approach</surname>
          </string-name>
          <string-name>
            <surname>”</surname>
          </string-name>
          , Pearson,
          <year>2013</year>
          .
        </mixed-citation>
      </ref>
      <ref id="ref17">
        <mixed-citation>
          [17]
          <string-name>
            <surname>Kurniawan</surname>
            ,
            <given-names>Y.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Gunawan</surname>
            ,
            <given-names>A.</given-names>
          </string-name>
          , and
          <string-name>
            <surname>Kurnia</surname>
            ,
            <given-names>S. G.</given-names>
          </string-name>
          ,
          <article-title>“Application of Business Intelligence to Support Marketing Strategies: A case study approach”</article-title>
          ,
          <source>Journal of Theoretical &amp; Applied Information Technology, 64.1</source>
          (
          <year>2014</year>
          ).
        </mixed-citation>
      </ref>
      <ref id="ref18">
        <mixed-citation>
          [18]
          <string-name>
            <surname>Chen</surname>
            ,
            <given-names>H.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Chiang</surname>
            ,
            <given-names>R. H.</given-names>
          </string-name>
          , and
          <string-name>
            <surname>Storey</surname>
            ,
            <given-names>V. C.</given-names>
          </string-name>
          , “
          <article-title>Business intelligence and analytics: From big data to big impact”</article-title>
          ,
          <source>MIS quarterly 36.4</source>
          (
          <year>2012</year>
          ):
          <fpage>1165</fpage>
          -
          <lpage>1188</lpage>
          .
        </mixed-citation>
      </ref>
      <ref id="ref19">
        <mixed-citation>
          [19]
          <string-name>
            <surname>Sun</surname>
            ,
            <given-names>Z.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Sun</surname>
            ,
            <given-names>L.</given-names>
          </string-name>
          and
          <string-name>
            <surname>Strang</surname>
            ,
            <given-names>K.</given-names>
          </string-name>
          , “
          <article-title>Big Data Analytics Services for Enhancing Business Intelligence”</article-title>
          ,
          <source>Journal of Computer Information Systems</source>
          ,
          <volume>58</volume>
          .2 (
          <year>2018</year>
          ):
          <fpage>162</fpage>
          -
          <lpage>169</lpage>
          .
        </mixed-citation>
      </ref>
      <ref id="ref20">
        <mixed-citation>
          [20]
          <string-name>
            <surname>Schulz</surname>
            ,
            <given-names>M.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Neuhaus</surname>
            ,
            <given-names>U.</given-names>
          </string-name>
          , Kaufmann, J.,
          <string-name>
            <surname>Kühnel</surname>
            ,
            <given-names>S.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Alekozai</surname>
            ,
            <given-names>E.M.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Rohde</surname>
            ,
            <given-names>H.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Badura</surname>
            ,
            <given-names>D.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Kerzel</surname>
            ,
            <given-names>U.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Lanquillon</surname>
            ,
            <given-names>C.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Daurer</surname>
            ,
            <given-names>S.</given-names>
          </string-name>
          <string-name>
            <surname>Günther</surname>
            ,
            <given-names>M.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Huber</surname>
            ,
            <given-names>L.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Thiée</surname>
            , L.-W., zur Heiden,
            <given-names>P.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Passlick</surname>
            ,
            <given-names>J.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Dieckmann</surname>
            ,
            <given-names>J.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Schwade</surname>
            ,
            <given-names>F.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Seyffarth</surname>
            ,
            <given-names>T.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Badewitz</surname>
            ,
            <given-names>W.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Rissler</surname>
            ,
            <given-names>R.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Sackmann</surname>
            ,
            <given-names>S.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Gölzer</surname>
            ,
            <given-names>P.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Welter</surname>
            ,
            <given-names>F.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Röth</surname>
            ,
            <given-names>J.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Seidelmann</surname>
            ,
            <given-names>J.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>Haneke</surname>
            ,
            <given-names>U.</given-names>
          </string-name>
          ,
          <string-name>
            <surname>„</surname>
            <given-names>DASC-PM</given-names>
          </string-name>
          <year>v1</year>
          .
          <article-title>1 - A Process Model for Data Science Projects”</article-title>
          , Elmshorn,
          <year>2022</year>
          .
        </mixed-citation>
      </ref>
      <ref id="ref21">
        <mixed-citation>
          [21]
          <string-name>
            <surname>Virkus</surname>
            ,
            <given-names>S.</given-names>
          </string-name>
          and
          <string-name>
            <surname>Garoufallou</surname>
          </string-name>
          , E.,
          <article-title>“Data Science from a Perspective of Computer Science”</article-title>
          ,
          <source>in Research Conference on Metadata and Semantics Research</source>
          (pp.
          <fpage>209</fpage>
          -
          <lpage>219</lpage>
          ). Springer, Cham,
          <year>2019</year>
          .
        </mixed-citation>
      </ref>
      <ref id="ref22">
        <mixed-citation>
          [22]
          <string-name>
            <surname>Taddy</surname>
            ,
            <given-names>M.</given-names>
          </string-name>
          , “
          <article-title>Business Data Science: Combining Machine Learning</article-title>
          and Economics to Optimize, Automate, and Accelerate Business Decisions”,
          <string-name>
            <surname>McGraw-Hill</surname>
            <given-names>Education</given-names>
          </string-name>
          , New York ,
          <year>2019</year>
          .
        </mixed-citation>
      </ref>
    </ref-list>
  </back>
</article>