Brunobkr/llama.cpp_AlgMor24_github
ΩFFFΣLLIa • llama.cpp • AlgMor24 ██████╗ ███████╗███████╗███████╗██╗ ██╗ ██╗ █████╗ ██╔═══██╗██╔════╝██╔════╝██╔════╝██║ ██║ ██║██╔══██╗ ██║ ██║█████╗ █████╗ █████╗ ██║ ██║ ██║███████║ ██║ ██║██╔══╝ ██╔══╝ ██╔══╝ ██║ ██║ ██║██╔══██║ ╚██████╔╝██║ ██║ ███████╗███████╗███████╗██║██║ ██║ ╚═════╝ ╚═╝ ╚═╝ ╚══════╝╚══════╝╚══════╝╚═╝╚═╝ ╚═╝ High-Performance LLM / VLM Inference & Autonomous Agentic Ecosystem… See the full description on the dataset page: https://huggingface.co/datasets/Brunobkr/llama.cpp_AlgMor24_github.
03k
1name: Server2 3on:4 workflow_dispatch: # allows manual triggering5 inputs:6 sha:7 description: 'Commit SHA1 to build'8 required: false9 type: string10 slow_tests:11 description: 'Run slow tests'12 required: true13 type: boolean14 push:15 branches:16 - master17 paths: [18 '.github/workflows/server.yml',19 '**/CMakeLists.txt',20 '**/Makefile',21 '**/*.h',22 '**/*.hpp',23 '**/*.c',24 '**/*.cpp',25 '**/*.cu',26 '**/*.swift',27 '**/*.m',28 'tools/server/**.*'29 ]30 pull_request:31 types: [opened, synchronize, reopened]32 paths: [33 '.github/workflows/server.yml',34 '**/CMakeLists.txt',35 '**/Makefile',36 '**/*.h',37 '**/*.hpp',38 '**/*.c',39 '**/*.cpp',40 '**/*.cu',41 '**/*.swift',42 '**/*.m',43 'tools/server/**.*'44 ]45 46env:47 LLAMA_ARG_LOG_COLORS: 148 LLAMA_ARG_LOG_PREFIX: 149 LLAMA_ARG_LOG_TIMESTAMPS: 150 LLAMA_ARG_LOG_VERBOSITY: 1051 52concurrency:53 group: ${{ github.workflow }}-${{ github.ref }}-${{ github.head_ref || github.run_id }}54 cancel-in-progress: true55 56jobs:57 ubuntu:58 runs-on: ubuntu-24.04-arm59 60 steps:61 - name: Dependencies62 id: depends63 run: |64 sudo apt-get update65 sudo apt-get -y install \66 build-essential \67 xxd \68 git \69 cmake \70 curl \71 wget \72 language-pack-en \73 libssl-dev74 75 - name: Clone76 id: checkout77 uses: actions/checkout@v678 with:79 fetch-depth: 080 ref: ${{ github.event.inputs.sha || github.event.pull_request.head.sha || github.sha || github.head_ref || github.ref_name }}81 82 - name: ccache83 uses: ggml-org/ccache-action@v1.2.2184 with:85 key: server-ubuntu-24.04-arm86 evict-old-files: 1d87 save: ${{ github.event_name == 'push' && github.ref == 'refs/heads/master' }}88 89 - name: Build90 id: cmake_build91 run: |92 cmake -B build \93 -DGGML_SCHED_NO_REALLOC=ON94 cmake --build build --config Release -j $(nproc) --target llama-server95 96 - name: Python setup97 id: setup_python98 uses: actions/setup-python@v699 with:100 python-version: '3.11'101 pip-install: -r tools/server/tests/requirements.txt102 103 - name: Tests104 id: server_integration_tests105 run: |106 cd tools/server/tests107 pytest -v -x -m "not slow"108 109 - name: Slow tests110 id: server_integration_tests_slow111 if: ${{ github.event.schedule || github.event.inputs.slow_tests == 'true' }}112 run: |113 cd tools/server/tests114 SLOW_TESTS=1 pytest -v -x115 116 - name: Tests (Backend sampling)117 id: server_integration_tests_backend_sampling118 run: |119 cd tools/server/tests120 export LLAMA_ARG_BACKEND_SAMPLING=1121 pytest -v -x -m "not slow"122 123 - name: Slow tests (Backend sampling)124 id: server_integration_tests_slow_backend_sampling125 if: ${{ github.event.schedule || github.event.inputs.slow_tests == 'true' }}126 run: |127 cd tools/server/tests128 export LLAMA_ARG_BACKEND_SAMPLING=1129 SLOW_TESTS=1 pytest -v -x130 131 windows:132 runs-on: windows-2025133 134 steps:135 - name: Clone136 id: checkout137 uses: actions/checkout@v6138 with:139 fetch-depth: 0140 ref: ${{ github.event.inputs.sha || github.event.pull_request.head.sha || github.sha || github.head_ref || github.ref_name }}141 142 - name: ccache143 uses: ggml-org/ccache-action@v1.2.21144 with:145 key: server-windows-2025-x64146 evict-old-files: 1d147 save: ${{ github.event_name == 'push' && github.ref == 'refs/heads/master' }}148 149 - name: Build150 id: cmake_build151 shell: cmd152 run: |153 cmake -B build -G "Ninja Multi-Config" ^154 -DCMAKE_TOOLCHAIN_FILE=cmake/x64-windows-llvm.cmake ^155 -DCMAKE_BUILD_TYPE=Release ^156 -DLLAMA_BUILD_BORINGSSL=ON ^157 -DGGML_SCHED_NO_REALLOC=ON158 set /A NINJA_JOBS=%NUMBER_OF_PROCESSORS%-1159 cmake --build build --config Release -j %NINJA_JOBS% --target llama-server160 161 - name: Python setup162 id: setup_python163 uses: actions/setup-python@v6164 with:165 python-version: '3.11'166 pip-install: -r tools/server/tests/requirements.txt167 168 - name: Tests169 id: server_integration_tests170 run: |171 cd tools/server/tests172 $env:PYTHONIOENCODING = ":replace"173 pytest -v -x -m "not slow"174 175 - name: Slow tests176 id: server_integration_tests_slow177 if: ${{ github.event.schedule || github.event.inputs.slow_tests == 'true' }}178 run: |179 cd tools/server/tests180 $env:SLOW_TESTS = "1"181 pytest -v -x182 