-
Notifications
You must be signed in to change notification settings - Fork 0
/
Copy pathDockerfile
78 lines (51 loc) · 1.34 KB
/
Dockerfile
1
2
3
4
5
6
7
8
9
10
11
12
13
14
15
16
17
18
19
20
21
22
23
24
25
26
27
28
29
30
31
32
33
34
35
36
37
38
39
40
41
42
43
44
45
46
47
48
49
50
51
52
53
54
55
56
57
58
59
60
61
62
63
64
65
66
67
68
69
70
71
72
73
74
75
76
77
78
FROM eclipse-temurin:21-jdk AS app_stage
WORKDIR /app
COPY build.gradle.kts .
COPY settings.gradle.kts .
COPY gradlew .
COPY gradle/ ./gradle
RUN ./gradlew clean
COPY src/ ./src/
RUN ./gradlew clean bootJar
FROM ubuntu:22.04 AS tesseract
LABEL authors="nenadjakic"
ARG DEBIAN_FRONTEND=noninteractive
RUN apt-get -y update
RUN apt-get -y install \
automake \
ca-certificates \
g++ \
git \
libtool \
libleptonica-dev \
make \
pkg-config
RUN apt-get -y install --no-install-recommends \
asciidoc \
docbook-xsl \
xsltproc
WORKDIR /src
ARG TESSERACT_VERSION
RUN git clone -b $TESSERACT_VERSION https://github.com/tesseract-ocr/tesseract.git
WORKDIR /src/tesseract
RUN ./autogen.sh
RUN ./configure
RUN make
RUN make install
RUN ldconfig
RUN apt-get -y install \
wget
WORKDIR /usr/local/share/tessdata/
COPY .tesseract/get-languages.sh .
COPY .tesseract/languages.txt .
RUN chmod +x ./get-languages.sh
RUN ./get-languages.sh
WORKDIR /usr/local/share/tessdata/scripts
COPY .tesseract/get-scripts.sh .
COPY .tesseract/scripts.txt .
RUN chmod +x ./get-scripts.sh
RUN ./get-scripts.sh
RUN apt-get -y install openjdk-21-jdk
WORKDIR /app
COPY --from=app_stage /app/build/libs/ocr-studio.jar .
CMD ["java", "-jar", "ocr-studio.jar"]