| |
|
| | --- |
| | language: |
| | - pt |
| | - en |
| | tags: |
| | - aes |
| | datasets: |
| | - kamel-usp/aes_enem_dataset |
| | base_model: microsoft/Phi-3.5-mini-instruct |
| | metrics: |
| | - accuracy |
| | - qwk |
| | library_name: peft |
| | model-index: |
| | - name: phi35-balanced-C4 |
| | results: |
| | - task: |
| | type: text-classification |
| | name: Automated Essay Score |
| | dataset: |
| | name: Automated Essay Score ENEM Dataset |
| | type: kamel-usp/aes_enem_dataset |
| | config: JBCS2025 |
| | split: test |
| | metrics: |
| | - name: Macro F1 |
| | type: f1 |
| | value: 0.29404571649185807 |
| | - name: QWK |
| | type: qwk |
| | value: 0.5570197668525089 |
| | - name: Weighted Macro F1 |
| | type: f1 |
| | value: 0.5930707679874385 |
| | --- |
| | # Model ID: phi35-balanced-C4 |
| | ## Results |
| | | | test_data | |
| | |:-----------------|------------:| |
| | | eval_accuracy | 0.572464 | |
| | | eval_RMSE | 29.6843 | |
| | | eval_QWK | 0.55702 | |
| | | eval_Macro_F1 | 0.294046 | |
| | | eval_Weighted_F1 | 0.593071 | |
| | | eval_Micro_F1 | 0.572464 | |
| | | eval_HDIV | 0.00724638 | |
| | |