chore(config): update environment and Docker configurations

- Removed default FIRECRAWL_API_URL from `.env.example` and `docker-compose.yml`.
- Added new provider configuration files to the Dockerfile, including hcnsec, novita, and vllm-mixed.
- Deleted the obsolete glm_flash_bedrock.yml configuration file.
This commit is contained in:
Dmitry Ng
2026-07-23 18:44:03 +03:00
parent deaff9eafe
commit 38c1e261a2
4 changed files with 9 additions and 183 deletions
+1 -1
View File
@@ -195,7 +195,7 @@ TAVILY_API_KEY=
## Firecrawl search engine API
## FIRECRAWL_API_URL is optional; leave the default for the cloud API or point it at a self-hosted instance
FIRECRAWL_API_KEY=
FIRECRAWL_API_URL=https://api.firecrawl.dev
FIRECRAWL_API_URL=
## Perplexity search engine API
PERPLEXITY_API_KEY=
+7 -1
View File
@@ -159,18 +159,24 @@ COPY --from=api-builder /licenses/backend /opt/pentagi/licenses/backend
COPY --from=frontend-compiler /licenses/frontend /opt/pentagi/licenses/frontend
# Copy provider configuration files
COPY examples/configs/atlas.provider.yml /opt/pentagi/conf/
COPY examples/configs/azure-openai.provider.yml /opt/pentagi/conf/
COPY examples/configs/bedrock-glm-flash.provider.yml /opt/pentagi/conf/
COPY examples/configs/custom-openai.provider.yml /opt/pentagi/conf/
COPY examples/configs/deepinfra.provider.yml /opt/pentagi/conf/
COPY examples/configs/deepseek.provider.yml /opt/pentagi/conf/
COPY examples/configs/hcnsec.provider.yml /opt/pentagi/conf/
COPY examples/configs/moonshot.provider.yml /opt/pentagi/conf/
COPY examples/configs/novita.provider.yml /opt/pentagi/conf/
COPY examples/configs/nvidia-glm-5.1.provider.yml /opt/pentagi/conf/
COPY examples/configs/ollama-cloud.provider.yml /opt/pentagi/conf/
COPY examples/configs/ollama-llama318b-instruct.provider.yml /opt/pentagi/conf/
COPY examples/configs/ollama-llama318b.provider.yml /opt/pentagi/conf/
COPY examples/configs/ollama-qwen332b-fp16-tc.provider.yml /opt/pentagi/conf/
COPY examples/configs/ollama-qwq32b-fp16-tc.provider.yml /opt/pentagi/conf/
COPY examples/configs/openrouter.provider.yml /opt/pentagi/conf/
COPY examples/configs/novita.provider.yml /opt/pentagi/conf/
COPY examples/configs/orcarouter.provider.yml /opt/pentagi/conf/
COPY examples/configs/vllm-mixed.provider.yml /opt/pentagi/conf/
COPY examples/configs/vllm-qwen3.5-27b-fp8-no-think.provider.yml /opt/pentagi/conf/
COPY examples/configs/vllm-qwen3.5-27b-fp8.provider.yml /opt/pentagi/conf/
COPY examples/configs/vllm-qwen3.6-27b-fp8-no-think.provider.yml /opt/pentagi/conf/
+1 -1
View File
@@ -156,7 +156,7 @@ services:
- TRAVERSAAL_API_KEY=${TRAVERSAAL_API_KEY:-}
- TAVILY_API_KEY=${TAVILY_API_KEY:-}
- FIRECRAWL_API_KEY=${FIRECRAWL_API_KEY:-}
- FIRECRAWL_API_URL=${FIRECRAWL_API_URL:-https://api.firecrawl.dev}
- FIRECRAWL_API_URL=${FIRECRAWL_API_URL:-}
- PERPLEXITY_API_KEY=${PERPLEXITY_API_KEY:-}
- PERPLEXITY_MODEL=${PERPLEXITY_MODEL:-sonar}
- PERPLEXITY_CONTEXT_SIZE=${PERPLEXITY_CONTEXT_SIZE:-low}
-180
View File
@@ -1,180 +0,0 @@
# GLM Flash configuration for AWS Bedrock (zai.glm-4.7-flash)
# Exported from the glm_flash provider stored in the database.
# To use, set BEDROCK_CONFIG_PATH=./glm_flash_bedrock.yml in your .env file.
name: glm_flash
simple:
model: zai.glm-4.7-flash
temperature: 0.5
top_p: 0.5
n: 1
max_tokens: 6000
price:
input: 0.15
output: 0.6
cache_read: 0
cache_write: 0
simple_json:
model: zai.glm-4.7-flash
temperature: 0.5
top_p: 0.5
n: 1
max_tokens: 63998
json: true
price:
input: 0.15
output: 0.6
cache_read: 0
cache_write: 0
primary_agent:
model: zai.glm-4.7-flash
temperature: 1.0
n: 1
max_tokens: 63999
reasoning:
effort: high
max_tokens: 2048
price:
input: 3.0
output: 15.0
cache_read: 0.3
cache_write: 3.75
assistant:
model: zai.glm-4.7-flash
temperature: 1.0
n: 1
max_tokens: 63998
reasoning:
effort: low
max_tokens: 1024
price:
input: 3.0
output: 15.0
cache_read: 0.3
cache_write: 3.75
generator:
model: zai.glm-4.7-flash
temperature: 1.0
n: 1
max_tokens: 63998
reasoning:
max_tokens: 4096
price:
input: 3.0
output: 15.0
cache_read: 0.3
cache_write: 3.75
refiner:
model: zai.glm-4.7-flash
temperature: 1.0
n: 1
max_tokens: 63998
reasoning:
effort: medium
max_tokens: 2048
price:
input: 3.0
output: 15.0
cache_read: 0.3
cache_write: 3.75
adviser:
model: zai.glm-4.7-flash
temperature: 1.0
n: 1
max_tokens: 63998
reasoning:
effort: high
max_tokens: 64000
price:
input: 5.0
output: 25.0
cache_read: 0.5
cache_write: 6.25
reflector:
model: zai.glm-4.7-flash
temperature: 1.0
n: 1
max_tokens: 4096
reasoning:
effort: low
max_tokens: 1024
price:
input: 1.0
output: 5.0
cache_read: 0.1
cache_write: 1.25
searcher:
model: zai.glm-4.7-flash
temperature: 1.0
n: 1
max_tokens: 8192
reasoning:
max_tokens: 1024
price:
input: 1.0
output: 5.0
cache_read: 0.1
cache_write: 1.25
enricher:
model: zai.glm-4.7-flash
temperature: 1.0
n: 1
max_tokens: 63998
reasoning:
max_tokens: 1024
price:
input: 1.0
output: 5.0
cache_read: 0.1
cache_write: 1.25
coder:
model: zai.glm-4.7-flash
temperature: 1.0
n: 1
max_tokens: 63997
reasoning:
effort: high
max_tokens: 2048
price:
input: 3.0
output: 15.0
cache_read: 0.3
cache_write: 3.75
installer:
model: zai.glm-4.7-flash
temperature: 1.0
n: 1
max_tokens: 63996
reasoning:
max_tokens: 1024
price:
input: 3.0
output: 15.0
cache_read: 0.3
cache_write: 3.75
pentester:
model: zai.glm-4.7-flash
temperature: 1.0
n: 1
max_tokens: 128000
reasoning:
effort: high
max_tokens: 1024
price:
input: 3.0
output: 15.0
cache_read: 0.3
cache_write: 3.75