{
  "model_details": {
    "provider": "e-infra",
    "model_requested": "qwen3.8-27b",
    "model_for_requests": "qwen3.8-27b",
    "api_base_url": "https://llm.ai.e-infra.cz/v1",
    "chat_completions_endpoint": "https://llm.ai.e-infra.cz/v1/chat/completions",
    "model_metadata": {
      "model_name": "qwen3.8-27b",
      "litellm_model": "qwen3.8-27b",
      "context_size": 262144.0,
      "max_output_tokens": 32000.0,
      "model_source": "https://huggingface.co/Qwen/Qwen3.8-27B-FP8",
      "quantization": "fp8",
      "capabilities": [
        "chat",
        "multimodal",
        "tools"
      ],
      "extra_body": {
        "max_tokens": 32000,
        "max_completion_tokens": 32000
      },
      "input_cost_per_token": 4.0000000000000003e-07,
      "output_cost_per_token": 3e-06
    },
    "model_source": "https://huggingface.co/Qwen/Qwen3.8-27B-FP8",
    "quantization": "fp8",
    "context_size": 262144.0,
    "max_output_tokens": 32000.0,
    "input_cost_per_token": 4.0000000000000003e-07,
    "output_cost_per_token": 3e-06
  },
  "run_config": {
    "input": [
      "data/input/prepositions.csv"
    ],
    "labels": "",
    "task_name": "Err. correct. prepositions",
    "task_description": "Correcting errors in prepositions in a transcription of learners' dialogue.",
    "tags": "error;correction;preposition;English",
    "model": "qwen3.8-27b",
    "temperature": null,
    "top_p": null,
    "top_k": null,
    "service_tier": "standard",
    "verbosity": null,
    "reasoning_effort": null,
    "thinking_level": null,
    "effort": null,
    "strict_control_acceptance": true,
    "provider": "e-infra",
    "system_prompt": "You are a meticulous annotator focusing on grammatical errors in the use of prepositions.\nYou will receive a short excerpt of a transcribed dialogue where the node word is a preposition. \nDecide whether it is a correct preposition given the context. \nReturn the correct preposition (even if it already is correct). If the preposition should be omitted (i.e. the grammatical correct solution is leaving it out), return 0.",
    "system_prompt_b64": null,
    "few_shot_examples": 0,
    "prompt_layout": "standard",
    "prompt_batch_size": 0,
    "cache_pad_target_tokens": 0,
    "prompt_cache_key": null,
    "openai_cache_breakpoint": false,
    "openrouter_cache_control": false,
    "cache_warmup_delay_seconds": 5.0,
    "gemini_cached_content": null,
    "requesty_auto_cache": null,
    "vertex_auto_adc_login": null,
    "vertex_access_token_refresh_seconds": null,
    "create_gemini_cache": false,
    "gemini_cache_ttl": 3600,
    "gemini_cache_ttl_autoupdate": true,
    "keep_gemini_cache": false,
    "enable_cot": true,
    "no_explanation": false,
    "logprobs": true,
    "calibration": true,
    "confusion_heatmap": true,
    "repeat_unclassified": true,
    "api_key_var": "E-INFRA_API_KEY",
    "api_base_var": "E-INFRA_BASE_URL",
    "max_retries": 3,
    "retry_delay": 5.0,
    "request_interval_ms": 0,
    "request_timeout_seconds": 30.0,
    "threads": 4,
    "prompt_log_detail": "full",
    "flush_rows": 100,
    "flush_seconds": 2.0,
    "validator_cmd": null,
    "validator_args": "",
    "validator_timeout": 5.0,
    "validator_prompt_max_candidates": 50,
    "validator_prompt_max_chars": 8000,
    "validator_exhausted_policy": "accept_blank_confidence",
    "validator_debug": false,
    "log_level": "INFO",
    "load_params": "data/metrics/prepositions__einfra__kimik3__2026-08-25-22-34__metrics.json"
  },
  "source_input_csv": "data/input/prepositions.csv",
  "source_output_csv": "data/output/prepositions__einfra__qwen3827b__2026-09-07-22-07.csv",
  "source_labels_csv": "",
  "cache_padding": {
    "enabled": false,
    "target_shared_prefix_tokens": 0,
    "calibration_shared_prefix_tokens": null,
    "target_prompt_tokens": 0,
    "calibration_prompt_tokens": null,
    "calibration_example_id": null,
    "applied_padding_tokens_estimate": 0,
    "examples_with_padding_applied": 0
  },
  "cache_warmup": {
    "enabled": false,
    "delay_seconds": 5.0,
    "attempted": false,
    "wait_applied": false,
    "reported_cache_write_tokens": 0
  },
  "request_control_summary": {
    "configured": {},
    "attempts_total": 1549,
    "attempts_with_control_telemetry": 1290,
    "per_control": {
      "reasoning_effort": {
        "configured_value": null,
        "requested_attempts": 0,
        "sent_attempts": 0,
        "accepted_attempts": 0,
        "rejected_attempts": 0,
        "missing_from_final_request_attempts": 0,
        "acceptance_rate": null,
        "rejected_reasons": {},
        "rejected_example_ids": []
      },
      "thinking_level": {
        "configured_value": null,
        "requested_attempts": 0,
        "sent_attempts": 0,
        "accepted_attempts": 0,
        "rejected_attempts": 0,
        "missing_from_final_request_attempts": 0,
        "acceptance_rate": null,
        "rejected_reasons": {},
        "rejected_example_ids": []
      },
      "effort": {
        "configured_value": null,
        "requested_attempts": 0,
        "sent_attempts": 0,
        "accepted_attempts": 0,
        "rejected_attempts": 0,
        "missing_from_final_request_attempts": 0,
        "acceptance_rate": null,
        "rejected_reasons": {},
        "rejected_example_ids": []
      },
      "verbosity": {
        "configured_value": null,
        "requested_attempts": 0,
        "sent_attempts": 0,
        "accepted_attempts": 0,
        "rejected_attempts": 0,
        "missing_from_final_request_attempts": 0,
        "acceptance_rate": null,
        "rejected_reasons": {},
        "rejected_example_ids": []
      },
      "prompt_cache_key": {
        "configured_value": null,
        "requested_attempts": 0,
        "sent_attempts": 0,
        "accepted_attempts": 0,
        "rejected_attempts": 0,
        "missing_from_final_request_attempts": 0,
        "acceptance_rate": null,
        "rejected_reasons": {},
        "rejected_example_ids": []
      },
      "openai_cache_breakpoint": {
        "configured_value": null,
        "requested_attempts": 0,
        "sent_attempts": 0,
        "accepted_attempts": 0,
        "rejected_attempts": 0,
        "missing_from_final_request_attempts": 0,
        "acceptance_rate": null,
        "rejected_reasons": {},
        "rejected_example_ids": []
      },
      "openrouter_cache_control": {
        "configured_value": null,
        "requested_attempts": 0,
        "sent_attempts": 0,
        "accepted_attempts": 0,
        "rejected_attempts": 0,
        "missing_from_final_request_attempts": 0,
        "acceptance_rate": null,
        "rejected_reasons": {},
        "rejected_example_ids": []
      },
      "gemini_cached_content": {
        "configured_value": null,
        "requested_attempts": 0,
        "sent_attempts": 0,
        "accepted_attempts": 0,
        "rejected_attempts": 0,
        "missing_from_final_request_attempts": 0,
        "acceptance_rate": null,
        "rejected_reasons": {},
        "rejected_example_ids": []
      },
      "requesty_auto_cache": {
        "configured_value": null,
        "requested_attempts": 0,
        "sent_attempts": 0,
        "accepted_attempts": 0,
        "rejected_attempts": 0,
        "missing_from_final_request_attempts": 0,
        "acceptance_rate": null,
        "rejected_reasons": {},
        "rejected_example_ids": []
      }
    }
  },
  "usage_metadata_summary": {
    "attempts_total": 1549,
    "attempts_with_usage_metadata": 1290,
    "attempts_with_cached_token_signals": 1287,
    "cached_tokens_total_estimate": 583360,
    "cache_read_tokens_total": 583360,
    "cache_write_tokens_total": 0,
    "cache_token_fields_totals": {
      "usage.prompt_tokens_details.cached_tokens": 583360
    },
    "attempts_with_gemini_cached_content_token_signals": 0,
    "gemini_cached_content_token_count_total": 0,
    "gemini_cached_content_token_fields_totals": {}
  },
  "token_usage_totals": {
    "attempts_total": 1549,
    "attempts_with_token_usage": 1290,
    "attempts_with_output_tokens": 1290,
    "attempts_with_cached_input_tokens": 1287,
    "attempts_with_thinking_tokens": 1290,
    "input_tokens_total": 854990,
    "cached_input_tokens_total": 583360,
    "non_cached_input_tokens_total": 271630,
    "output_tokens_total": 1156562,
    "thinking_tokens_total": 1044093,
    "output_tokens_definition": "total_tokens - prompt_tokens (or completion_tokens + thinking_tokens fallback)"
  },
  "truth_label_count": 1283,
  "prediction_count": 1283,
  "evaluated_example_count": 1283,
  "calibration_metrics": {
    "available": true,
    "sample_count": 1281,
    "bin_count": 10,
    "ece": 0.018969555035134346,
    "mce": 0.42000000000000004,
    "brier_score": 0.08331506635441058
  },
  "first_prompt_timestamp": "2026-09-07T20:07:25.632826Z",
  "last_prompt_timestamp": "2026-09-07T20:45:24.360686Z",
  "overall_time_seconds": 2263.2552490000003,
  "overall_time_human": "37m 43s",
  "accuracy": 0.8994544037412315,
  "cohen_kappa": 0.781584743514793,
  "macro_precision": 0.44089289274322174,
  "macro_recall": 0.4111106540910435,
  "macro_f1": 0.4037956554221216,
  "per_label": {
    "0": {
      "precision": 0.2777777777777778,
      "recall": 0.6,
      "f1": 0.37974683544303794,
      "support": 25
    },
    "about": {
      "precision": 0.5,
      "recall": 0.5,
      "f1": 0.5,
      "support": 2
    },
    "around": {
      "precision": 0.0,
      "recall": 0.0,
      "f1": 0.0,
      "support": 1
    },
    "as": {
      "precision": 0.0,
      "recall": 0.0,
      "f1": 0.0,
      "support": 1
    },
    "at": {
      "precision": 0.9242424242424242,
      "recall": 0.9104477611940298,
      "f1": 0.9172932330827067,
      "support": 201
    },
    "behind": {
      "precision": 1.0,
      "recall": 1.0,
      "f1": 1.0,
      "support": 1
    },
    "during": {
      "precision": 1.0,
      "recall": 0.5,
      "f1": 0.6666666666666666,
      "support": 4
    },
    "for": {
      "precision": 0.0,
      "recall": 0.0,
      "f1": 0.0,
      "support": 0
    },
    "from": {
      "precision": 0.6666666666666666,
      "recall": 1.0,
      "f1": 0.8,
      "support": 2
    },
    "in": {
      "precision": 0.944078947368421,
      "recall": 0.9368879216539717,
      "f1": 0.9404696886947024,
      "support": 919
    },
    "into": {
      "precision": 0.0,
      "recall": 0.0,
      "f1": 0.0,
      "support": 1
    },
    "of": {
      "precision": 0.0,
      "recall": 0.0,
      "f1": 0.0,
      "support": 2
    },
    "on": {
      "precision": 0.8526315789473684,
      "recall": 0.8804347826086957,
      "f1": 0.8663101604278075,
      "support": 92
    },
    "since": {
      "precision": 0.0,
      "recall": 0.0,
      "f1": 0.0,
      "support": 0
    },
    "to": {
      "precision": 0.8888888888888888,
      "recall": 0.25,
      "f1": 0.3902439024390244,
      "support": 32
    },
    "with": {
      "precision": 0.0,
      "recall": 0.0,
      "f1": 0.0,
      "support": 0
    }
  },
  "labels": [
    "0",
    "about",
    "around",
    "as",
    "at",
    "behind",
    "during",
    "for",
    "from",
    "in",
    "into",
    "of",
    "on",
    "since",
    "to",
    "with"
  ],
  "label_count": 16,
  "total_examples": 1283,
  "confusion_matrix_sparse": [
    [
      0,
      0,
      15
    ],
    [
      0,
      4,
      1
    ],
    [
      0,
      9,
      8
    ],
    [
      0,
      12,
      1
    ],
    [
      1,
      1,
      1
    ],
    [
      1,
      9,
      1
    ],
    [
      2,
      9,
      1
    ],
    [
      3,
      9,
      1
    ],
    [
      4,
      0,
      1
    ],
    [
      4,
      4,
      183
    ],
    [
      4,
      9,
      12
    ],
    [
      4,
      12,
      5
    ],
    [
      5,
      5,
      1
    ],
    [
      6,
      6,
      2
    ],
    [
      6,
      9,
      2
    ],
    [
      8,
      8,
      2
    ],
    [
      9,
      0,
      32
    ],
    [
      9,
      1,
      1
    ],
    [
      9,
      4,
      13
    ],
    [
      9,
      7,
      2
    ],
    [
      9,
      8,
      1
    ],
    [
      9,
      9,
      861
    ],
    [
      9,
      11,
      1
    ],
    [
      9,
      12,
      7
    ],
    [
      9,
      14,
      1
    ],
    [
      10,
      9,
      1
    ],
    [
      11,
      0,
      1
    ],
    [
      11,
      9,
      1
    ],
    [
      12,
      0,
      4
    ],
    [
      12,
      9,
      6
    ],
    [
      12,
      12,
      81
    ],
    [
      12,
      13,
      1
    ],
    [
      14,
      0,
      1
    ],
    [
      14,
      4,
      1
    ],
    [
      14,
      9,
      18
    ],
    [
      14,
      10,
      1
    ],
    [
      14,
      12,
      1
    ],
    [
      14,
      14,
      8
    ],
    [
      14,
      15,
      2
    ]
  ],
  "confusion_matrix_format": "sparse_triplets",
  "confusion_matrix_nonzero_cells": 39,
  "label_metrics_available": true
}